{"id":"8a798746-e359-4cde-9bdc-603f88d86827","slug":"crewai-ismail-2001-agent-bench","name":"agent-bench","description":"Industrial-grade benchmarking engine for AI agents. Define test scenarios in YAML, run high-performance parallel evaluations, and generate premium glassmorphism reports. Supports LangGraph, CrewAI, AutoGen, and custom agent stacks.","canonicalUrl":"https://www.xpersona.co/agent/crewai-ismail-2001-agent-bench","sourceUrl":"https://github.com/Ismail-2001/agent-bench","homepage":null,"source":"GITHUB_OPENCLEW","vendor":{"slug":"ismail-2001","label":"Ismail 2001","url":"https://github.com/Ismail-2001/agent-bench"},"protocols":["OPENCLEW"],"capabilities":["crewai","multi-agent"],"trustScore":null,"trustConfidence":"unknown","artifactCount":0,"benchmarkCount":0,"lastRelease":null,"freshnessAt":"2026-05-18T06:45:23.700Z","freshnessLabel":"May 18, 2026","securityReviewed":true,"openapiReady":false,"stats":[{"label":"Trust score","value":"Unknown"},{"label":"Compatibility","value":"OpenClaw"},{"label":"Freshness","value":"May 18, 2026"},{"label":"Vendor","value":"Ismail 2001"},{"label":"Artifacts","value":"0"},{"label":"Benchmarks","value":"0"},{"label":"Last release","value":"Unpublished"}],"factsPreview":[{"factKey":"vendor","label":"Vendor","value":"Ismail 2001","category":"vendor","href":"https://github.com/Ismail-2001/agent-bench","sourceUrl":"https://github.com/Ismail-2001/agent-bench","sourceType":"profile","confidence":"medium","observedAt":"2026-05-18T06:45:23.701Z","isPublic":true,"metadata":{}},{"factKey":"protocols","label":"Protocol compatibility","value":"OpenClaw","category":"compatibility","href":"https://www.xpersona.co/api/v1/agents/crewai-ismail-2001-agent-bench/contract","sourceUrl":"https://www.xpersona.co/api/v1/agents/crewai-ismail-2001-agent-bench/contract","sourceType":"contract","confidence":"medium","observedAt":"2026-05-18T06:45:23.701Z","isPublic":true,"metadata":{}},{"factKey":"traction","label":"Adoption signal","value":"2 GitHub stars","category":"adoption","href":"https://github.com/Ismail-2001/agent-bench","sourceUrl":"https://github.com/Ismail-2001/agent-bench","sourceType":"profile","confidence":"medium","observedAt":"2026-05-18T06:45:23.701Z","isPublic":true,"metadata":{}},{"factKey":"docs_crawl","label":"Crawlable docs","value":"6 indexed pages on the official domain","category":"integration","href":"https://github.com/login?return_to=https%3A%2F%2Fgithub.com%2Fopenclaw%2Fskills%2Ftree%2Fmain%2Fskills%2Fasleep123%2Fcaldav-calendar","sourceUrl":"https://github.com/login?return_to=https%3A%2F%2Fgithub.com%2Fopenclaw%2Fskills%2Ftree%2Fmain%2Fskills%2Fasleep123%2Fcaldav-calendar","sourceType":"search_document","confidence":"medium","observedAt":"2026-04-15T05:03:46.393Z","isPublic":true,"metadata":{}},{"factKey":"handshake_status","label":"Handshake status","value":"UNKNOWN","category":"security","href":"https://www.xpersona.co/api/v1/agents/crewai-ismail-2001-agent-bench/trust","sourceUrl":"https://www.xpersona.co/api/v1/agents/crewai-ismail-2001-agent-bench/trust","sourceType":"trust","confidence":"medium","observedAt":null,"isPublic":true,"metadata":{}}],"highlights":["2 GitHub stars","Trust evidence available"],"agentCard":{"name":"agent-bench","description":"Industrial-grade benchmarking engine for AI agents. Define test scenarios in YAML, run high-performance parallel evaluations, and generate premium glassmorphism reports. Supports LangGraph, CrewAI, AutoGen, and custom agent stacks.","source":"GITHUB_OPENCLEW","sourceId":"crewai:1197011580","repository":"https://github.com/Ismail-2001/agent-bench","documentation":"https://www.xpersona.co/agent/crewai-ismail-2001-agent-bench","protocols":["OPENCLEW"],"capabilities":["crewai","multi-agent"],"languages":["python"],"install":{"command":"git clone https://github.com/Ismail-2001/agent-bench.git","ecosystem":"git"},"examples":[{"kind":"example","language":"bash","snippet":"pip install agentbench"},{"kind":"example","language":"yaml","snippet":"name: \"basic-research\"\ntasks:\n  - id: \"compare-frameworks\"\n    input: \"Compare LangGraph and CrewAI for production systems in 2026.\"\n    criteria:\n      - type: contains_all\n        values: [\"LangGraph\", \"CrewAI\"]\n      - type: min_length\n        value: 200\n      - type: llm_judge\n        prompt: \"Does this provide a technical comparison? Score 0-10.\"\n        threshold: 7\n    limits:\n      max_tokens: 50000\n      max_latency_seconds: 60"}]}}