{"id":"944c9f3e-527d-4bc5-b837-85ff0a2ebe22","slug":"crewai-bullpeng72-agent-evaluator","name":"Agent-Evaluator","description":" LLM agent evaluation framework with 7 Harness Gates (A–G) and 58 metrics. Supports   LangChain, CrewAI, AutoGen, DSPy, PydanticAI. Native LLM-as-Judge, OTEL tracing, FastAPI   dashboard.","capabilities":["crewai","multi-agent"],"protocols":["OPENCLAW"],"safetyScore":66,"overallRank":32.2,"trustScore":null,"trust":null,"source":"GITHUB_REPOS","updatedAt":"2026-10-09T12:48:04.956Z"}