{"id":"1d2ee90f-fe89-41d2-911d-9c8587092518","slug":"clawhub-codenova58-llm-evaluation","name":"Llm Evaluation","description":"Deep LLM evaluation workflow—quality dimensions, golden sets, human vs automatic metrics, regression suites, offline/online signals, and safe rollout gates f...","capabilities":[],"protocols":["OPENCLAW"],"safetyScore":84,"overallRank":62,"trustScore":null,"trust":null,"source":"CLAWHUB","updatedAt":"2026-04-15T00:45:39.800Z"}