{"id":"7d992139-125c-439a-a822-ef3602d6c30d","slug":"clawhub-globalcaos-memory-bench-pioneer","name":"TinkerClaw Memory Bench","description":"Be one of the first to benchmark your agent's memory — and help shape how AI remembers. Peer-review-grade evaluation (LLM-as-judge, nDCG/MAP/MRR with 95% CIs, ablations) against your live memory system. Runs entirely LOCALLY by default — no memory content leaves your machine, and excerpts are redacted even on the local path. The optional OpenAI judge is opt-in, prints the exact request body it would send, redacts secrets first, requires typed consent, and cannot be switched on by an unattended run. Submitting results is a separate confirmed step that validates the report against the full schema and previews every field in it, and identifies you only if you pass --contributor. Built for the TinkerClaw fork — github.com/globalcaos/tinkerclaw. See Permissions, Data Flow & Consent.","canonicalUrl":"https://www.xpersona.co/agent/clawhub-globalcaos-memory-bench-pioneer","sourceUrl":"https://clawhub.ai/globalcaos/memory-bench-pioneer","homepage":"https://clawhub.ai/globalcaos/skills/memory-bench-pioneer","source":"CLAWHUB","vendor":{"slug":"clawhub","label":"Clawhub","url":"https://clawhub.ai/globalcaos/skills/memory-bench-pioneer"},"protocols":["OPENCLEW"],"capabilities":[],"trustScore":null,"trustConfidence":"unknown","artifactCount":0,"benchmarkCount":0,"lastRelease":"2.1.3","freshnessAt":"2026-10-11T06:14:48.568Z","freshnessLabel":"Oct 11, 2026","securityReviewed":true,"openapiReady":false,"stats":[{"label":"Trust score","value":"Unknown"},{"label":"Compatibility","value":"OpenClaw"},{"label":"Freshness","value":"Oct 11, 2026"},{"label":"Vendor","value":"Clawhub"},{"label":"Artifacts","value":"0"},{"label":"Benchmarks","value":"0"},{"label":"Last release","value":"2.1.3"}],"factsPreview":[{"factKey":"vendor","category":"vendor","label":"Vendor","value":"Clawhub","href":"https://clawhub.ai/globalcaos/skills/memory-bench-pioneer","sourceUrl":"https://clawhub.ai/globalcaos/skills/memory-bench-pioneer","sourceType":"profile","confidence":"medium","observedAt":"2026-10-11T06:14:48.633Z","isPublic":true},{"factKey":"protocols","category":"compatibility","label":"Protocol compatibility","value":"OpenClaw","href":"https://www.xpersona.co/api/v1/agents/clawhub-globalcaos-memory-bench-pioneer/contract","sourceUrl":"https://www.xpersona.co/api/v1/agents/clawhub-globalcaos-memory-bench-pioneer/contract","sourceType":"contract","confidence":"medium","observedAt":"2026-10-11T06:14:48.633Z","isPublic":true},{"factKey":"traction","category":"adoption","label":"Adoption signal","value":"1.1K downloads","href":"https://clawhub.ai/globalcaos/memory-bench-pioneer","sourceUrl":"https://clawhub.ai/globalcaos/memory-bench-pioneer","sourceType":"profile","confidence":"medium","observedAt":"2026-10-11T06:14:48.633Z","isPublic":true},{"factKey":"latest_release","category":"release","label":"Latest release","value":"2.1.3","href":"https://clawhub.ai/globalcaos/memory-bench-pioneer","sourceUrl":"https://clawhub.ai/globalcaos/memory-bench-pioneer","sourceType":"release","confidence":"medium","observedAt":"2026-09-09T09:04:13.102Z","isPublic":true},{"factKey":"handshake_status","category":"security","label":"Handshake status","value":"UNKNOWN","href":"https://www.xpersona.co/api/v1/agents/clawhub-globalcaos-memory-bench-pioneer/trust","sourceUrl":"https://www.xpersona.co/api/v1/agents/clawhub-globalcaos-memory-bench-pioneer/trust","sourceType":"trust","confidence":"medium","observedAt":null,"isPublic":true}],"highlights":["1.1K downloads","Trust evidence available"],"agentCard":{"name":"TinkerClaw Memory Bench","description":"Be one of the first to benchmark your agent's memory — and help shape how AI remembers. Peer-review-grade evaluation (LLM-as-judge, nDCG/MAP/MRR with 95% CIs, ablations) against your live memory system. Runs entirely LOCALLY by default — no memory content leaves your machine, and excerpts are redacted even on the local path. The optional OpenAI judge is opt-in, prints the exact request body it would send, redacts secrets first, requires typed consent, and cannot be switched on by an unattended run. Submitting results is a separate confirmed step that validates the report against the full schema and previews every field in it, and identifies you only if you pass --contributor. Built for the TinkerClaw fork — github.com/globalcaos/tinkerclaw. See Permissions, Data Flow & Consent.","source":"CLAWHUB","sourceId":"clawhub:s17324vfeqe0ptzp84z2z9vttx883zdg:memory-bench-pioneer","homepage":"https://clawhub.ai/globalcaos/skills/memory-bench-pioneer","repository":"https://clawhub.ai/globalcaos/memory-bench-pioneer","documentation":"https://www.xpersona.co/agent/clawhub-globalcaos-memory-bench-pioneer","protocols":["OPENCLEW"],"examples":[{"kind":"example","language":"bash","snippet":"# Local judge — the default. Nothing leaves your machine.\npython3 scripts/rate.py --queries 30 --judge local --ablation\n\n# Stronger judge, but it TRANSMITS retrieved memory excerpts to OpenAI.\n# You will be asked to type 'send' before the first request.\npython3 scripts/rate.py --queries 30 --judge openai --ablation\n\n# Benchmark a copy instead of your live database\npython3 scripts/rate.py --db /tmp/memory-copy.db --judge local\n\n# Custom test set\npython3 scripts/rate.py --testset path/to/queries.json --judge local"},{"kind":"example","language":"bash","snippet":"# Anonymous — the default\npython3 scripts/collect.py --days 14 --output /tmp/memory-bench-report.json\n\n# Attributed to you (your username goes in the report, and it becomes public if you submit)\npython3 scripts/collect.py --days 14 --contributor YOUR_GITHUB_USER --output /tmp/memory-bench-report.json"}]}}