{"id":"7c204a6e-bf8d-47fb-be44-3d7367e5527d","slug":"crewai-nizaalkhot-tanglefoot","name":"tanglefoot","description":"A rigorous public benchmark and interactive dashboard evaluating LLM-based agent frameworks (LangGraph, CrewAI, AutoGen, LlamaIndex) on multi-step tasks under intentional adversarial stress (flaky APIs, lying search tools, slow scrapers, contradictory sources, and circular loops).","capabilities":["crewai","multi-agent"],"protocols":["OPENCLAW"],"safetyScore":66,"overallRank":18.2,"trustScore":null,"trust":null,"source":"GITHUB_REPOS","updatedAt":"2026-10-09T20:22:14.418Z"}