{"id":"f7c53faa-f0da-435f-8412-6480bdcf0850","entityType":"agent","slug":"clawhub-xiaoyaliu00-document-reader","name":"document-reader","canonicalUrl":"https://www.xpersona.co/agent/clawhub-xiaoyaliu00-document-reader","canonicalPath":"/agent/clawhub-xiaoyaliu00-document-reader","generatedAt":"2026-10-09T06:09:42.373Z","source":"CLAWHUB","claimStatus":"UNCLAIMED","verificationTier":"NONE","summary":{"evidence":{"source":"editorial-content","verified":true,"confidence":"high","updatedAt":"2026-10-09T05:56:35.991Z","emptyReason":null},"description":"通用文档读取工具，支持 PDF/DOCX/XLSX/PPTX/RTF/ODT 等多种文档格式，也支持 ZIP/TAR.GZ/RAR/7Z 等主流压缩包内文档直接读取 Skill: document-reader Owner: xiaoyaliu00 Summary: 通用文档读取工具，支持 PDF/DOCX/XLSX/PPTX/RTF/ODT 等多种文档格式，也支持 ZIP/TAR.GZ/RAR/7Z 等主流压缩包内文档直接读取 Tags: latest:1.0.0 Version history: v1.0.0 | 2026-04-08T08:42:48.169Z | user Initial release: support PDF/DOCX/XLSX/PPTX/RTF/ODT documents and ZIP/TAR/RAR/7Z compressed files Archive index: Archive v1.0.0: 4 files, 6842 bytes Files: scripts/document_reader.py (17331b), skill-card.md (1","descriptionLabel":"Technical summary","evidenceSummary":"Capability contract not published. No trust telemetry is available yet. 4.2K downloads reported by the source. Last updated 10/9/2026.","installCommand":"clawhub skill install s176v3w9m19hha3zy3k0nga2t184e5sa:document-reader","sourceUrl":"https://clawhub.ai/xiaoyaliu00/document-reader","homepage":"https://clawhub.ai/xiaoyaliu00/skills/document-reader","primaryLinks":[{"label":"View on ClawHub","url":"https://clawhub.ai/xiaoyaliu00/document-reader","kind":"source"},{"label":"Homepage","url":"https://clawhub.ai/xiaoyaliu00/skills/document-reader","kind":"homepage"}],"safetyScore":84,"overallRank":62,"popularityScore":41,"trustScore":null,"claimedByName":null,"isOwner":false,"seoDescription":"通用文档读取工具，支持 PDF/DOCX/XLSX/PPTX/RTF/ODT 等多种文档格式，也支持 ZIP/TAR.GZ/RAR/7Z 等主流压缩包内文档直接读取 Skill: document-reader Owner: xiaoyaliu00 Summary: 通用文档读取工具，支持 PDF/DOCX/XLSX/"},"coverage":{"evidence":{"source":"public-profile","verified":false,"confidence":"medium","updatedAt":"2026-10-09T05:56:35.991Z","emptyReason":null},"protocols":[{"protocol":"OPENCLEW","label":"OpenClaw","status":"self-declared","notes":"Declared in the public agent profile."}],"capabilities":[],"verifiedCount":0,"selfDeclaredCount":1,"capabilityMatrix":{"rows":[{"key":"OPENCLEW","type":"protocol","support":"unknown","confidenceSource":"profile","notes":"Listed on profile"}],"flattenedTokens":"protocol:OPENCLEW|unknown|profile"}},"adoption":{"evidence":{"source":"CLAWHUB","verified":false,"confidence":"medium","updatedAt":"2026-10-09T05:56:35.991Z","emptyReason":null},"stars":null,"forks":null,"downloads":4151,"packageName":null,"latestVersion":"1.0.0","tractionLabel":"4.2K downloads"},"release":{"evidence":{"source":"CLAWHUB","verified":false,"confidence":"medium","updatedAt":"2026-10-09T05:56:35.991Z","emptyReason":null},"lastUpdatedAt":"2026-10-09T05:56:35.991Z","lastCrawledAt":"2026-10-09T05:56:35.991Z","lastIndexedAt":null,"nextCrawlAt":"2026-10-10T05:56:35.991Z","lastVerifiedAt":null,"highlights":[{"version":"1.0.0","createdAt":"2026-04-08T08:42:48.169Z","changelog":"Initial release: support PDF/DOCX/XLSX/PPTX/RTF/ODT documents and ZIP/TAR/RAR/7Z compressed files","fileCount":4,"zipByteSize":6842}]},"execution":{"evidence":{"source":"CLAWHUB","verified":false,"confidence":"low","updatedAt":null,"emptyReason":"No published capability contract is available yet."},"installCommand":"clawhub skill install s176v3w9m19hha3zy3k0nga2t184e5sa:document-reader","setupComplexity":"low","setupSteps":["Setup complexity is LOW. This package is likely designed for quick installation with minimal external side-effects.","Final validation: Expose the agent to a mock request payload inside a sandbox and trace the network egress before allowing access to real customer data."],"contract":{"contractStatus":"missing","authModes":[],"requires":[],"forbidden":[],"supportsMcp":false,"supportsA2a":false,"supportsStreaming":false,"inputSchemaRef":null,"outputSchemaRef":null,"dataRegion":null,"contractUpdatedAt":null,"sourceUpdatedAt":null,"freshnessSeconds":null},"invocationGuide":{"preferredApi":{"snapshotUrl":"https://www.xpersona.co/api/v1/agents/clawhub-xiaoyaliu00-document-reader/snapshot","contractUrl":"https://www.xpersona.co/api/v1/agents/clawhub-xiaoyaliu00-document-reader/contract","trustUrl":"https://www.xpersona.co/api/v1/agents/clawhub-xiaoyaliu00-document-reader/trust"},"curlExamples":["curl -s \"https://www.xpersona.co/api/v1/agents/clawhub-xiaoyaliu00-document-reader/snapshot\"","curl -s \"https://www.xpersona.co/api/v1/agents/clawhub-xiaoyaliu00-document-reader/contract\"","curl -s \"https://www.xpersona.co/api/v1/agents/clawhub-xiaoyaliu00-document-reader/trust\""],"jsonRequestTemplate":{"query":"summarize this repo","constraints":{"maxLatencyMs":2000,"protocolPreference":["OPENCLEW"]}},"jsonResponseTemplate":{"ok":true,"result":{"summary":"...","confidence":0.9},"meta":{"source":"CLAWHUB","generatedAt":"2026-10-09T06:09:42.373Z"}},"retryPolicy":{"maxAttempts":3,"backoffMs":[500,1500,3500],"retryableConditions":["HTTP_429","HTTP_503","NETWORK_TIMEOUT"]}},"endpoints":{"dossierUrl":"https://www.xpersona.co/api/v1/agents/clawhub-xiaoyaliu00-document-reader/dossier","snapshotUrl":"https://www.xpersona.co/api/v1/agents/clawhub-xiaoyaliu00-document-reader/snapshot","contractUrl":"https://www.xpersona.co/api/v1/agents/clawhub-xiaoyaliu00-document-reader/contract","trustUrl":"https://www.xpersona.co/api/v1/agents/clawhub-xiaoyaliu00-document-reader/trust"}},"reliability":{"evidence":{"source":"runtime-metrics","verified":false,"confidence":"low","updatedAt":null,"emptyReason":"No trust, reliability, or runtime telemetry is available."},"trust":{"status":"unavailable","handshakeStatus":"UNKNOWN","verificationFreshnessHours":null,"reputationScore":null,"p95LatencyMs":null,"successRate30d":null,"fallbackRate":null,"attempts30d":null,"trustUpdatedAt":null,"trustConfidence":"unknown","sourceUpdatedAt":null,"freshnessSeconds":null},"decisionGuardrails":{"doNotUseIf":["Contract metadata is missing or unavailable for deterministic execution."],"safeUseWhen":[],"riskFlags":["missing_or_unavailable_contract","trust_data_unavailable","schema_references_missing"],"operationalConfidence":"low"},"executionMetrics":{"observedLatencyMsP50":null,"observedLatencyMsP95":null,"estimatedCostUsd":null,"uptime30d":null,"rateLimitRpm":null,"rateLimitBurst":null,"lastVerifiedAt":null,"verificationSource":null},"runtimeMetrics":{"successRate":null,"avgLatencyMs":null,"avgCostUsd":null,"hallucinationRate":null,"retryRate":null,"disputeRate":null,"p50Latency":null,"p95Latency":null,"lastUpdated":null}},"benchmarks":{"evidence":{"source":"no-benchmark-data","verified":false,"confidence":"low","updatedAt":null,"emptyReason":"No benchmark suites or observed failure patterns are available."},"suites":[],"failurePatterns":[]},"artifacts":{"evidence":{"source":"CLAWHUB","verified":false,"confidence":"high","updatedAt":"2026-10-09T05:56:35.991Z","emptyReason":null},"readme":"Skill: document-reader\n\nOwner: xiaoyaliu00\n\nSummary: 通用文档读取工具，支持 PDF/DOCX/XLSX/PPTX/RTF/ODT 等多种文档格式，也支持 ZIP/TAR.GZ/RAR/7Z 等主流压缩包内文档直接读取\n\nTags: latest:1.0.0\n\nVersion history:\n\nv1.0.0 | 2026-04-08T08:42:48.169Z | user\n\nInitial release: support PDF/DOCX/XLSX/PPTX/RTF/ODT documents and ZIP/TAR/RAR/7Z compressed files\n\nArchive index:\n\nArchive v1.0.0: 4 files, 6842 bytes\n\nFiles: scripts/document_reader.py (17331b), skill-card.md (1896b), SKILL.md (3690b), _meta.json (134b)\n\nFile v1.0.0:SKILL.md\n\n---\nname: document-reader\ndescription: \"通用文档读取工具，支持 PDF/DOCX/XLSX/PPTX/RTF/ODT 等多种文档格式，也支持 ZIP/TAR.GZ/RAR/7Z 等主流压缩包内文档直接读取\"\nversion: \"1.0.0\"\n---\n\n# Document Reader - 通用文档读取技能\n\n读取各种格式文档的内容，直接输出文本供 AI 分析。支持压缩包内文档直接读取，无需手动解压。\n\n## 📦 支持格式\n\n### 文档\n| 格式 | 扩展名 | 说明 |\n|------|--------|------|\n| PDF | `.pdf` | 支持文本提取，依赖 poppler-utils |\n| Microsoft Word | `.docx` | 完整提取所有段落文本 |\n| Microsoft Excel | `.xlsx` | 按 Sheet 输出，每个 Sheet 输出为表格文本 |\n| Microsoft PowerPoint | `.pptx` | 按 Slide 分块输出 |\n| Rich Text Format | `.rtf` |  |\n| OpenDocument Text | `.odt` |  |\n| HTML | `.html`/`.htm` | 提取正文文本 |\n| 纯文本 | `.txt`/`.md`/`.json`/`.xml`/`.py`/`.js` 等 | 直接读取 |\n\n### 压缩包\n| 格式 | 扩展名 | 功能 |\n|------|--------|------|\n| ZIP | `.zip` | 列出文件 ➜ 读取指定文档 |\n| TAR | `.tar`/`.tar.gz`/`.tgz`/`.tar.bz2` | 列出文件 ➜ 读取指定文档 |\n| RAR | `.rar` | 列出文件 ➜ 读取指定文档 |\n| 7-Zip | `.7z` | 列出文件 ➜ 读取指定文档 |\n\n## 🚀 快速开始\n\n### 依赖安装\n\n**Python 包：**\n```bash\npip install textract python-docx openpyxl python-pptx rarfile py7zr --break-system-packages\n```\n\n**系统依赖（Ubuntu/Debian）：**\n```bash\napt-get install -y poppler-utils antiword unrtf tidy libxml2-dev libxslt1-dev\n```\n\n## 💡 使用示例\n\n### 1. 直接读取本地文档\n\n```bash\n# 读取 PDF 文件\npython {baseDir}/scripts/document_reader.py --file /path/to/document.pdf\n\n# 读取 Word 文档\npython {baseDir}/scripts/document_reader.py --file /path/to/report.docx\n\n# 读取 Excel 文件（输出带 Sheet 分隔的表格文本）\npython {baseDir}/scripts/document_reader.py --file /path/to/data.xlsx\n\n# JSON 格式输出（方便程序处理）\npython {baseDir}/scripts/document_reader.py --file /path/to/data.xlsx --format json\n```\n\n### 2. 处理压缩包\n\n**先列出压缩包里有哪些文件：**\n```bash\n# 列出 ZIP 包内容\npython {baseDir}/scripts/document_reader.py --list /path/to/archive.zip\n\n# 列出 RAR 包内容\npython {baseDir}/scripts/document_reader.py --list /path/to/archive.rar\n\n# 列出 7z 包内容\npython {baseDir}/scripts/document_reader.py --list /path/to/archive.7z\n```\n\n**读取压缩包里的指定文档：**\n```bash\n# 读取 ZIP 包内的 Word 文档\npython {baseDir}/scripts/document_reader.py --file /path/to/archive.zip --inner-path document.docx\n\n# 读取 7z 包内的 PDF\npython {baseDir}/scripts/document_reader.py --file /path/to/archive.7z --inner-path report.pdf\n```\n\n### 输出示例\n\n**读取文档：**\n```\n=== report.pdf ===\n\n# 项目进度报告\n\n## 本周完成\n\n1. 完成了前端界面开发\n2. 修复了三个 Bug\n3. 编写了接口文档\n\n...\n```\n\n**列出压缩包：**\n```\nArchive: data.zip\nFound 3 file(s):\n\n  readme.txt\n  docs/report.pdf\n  data/sheet.xlsx\n```\n\n## ✨ 特性\n\n- 🎯 **开箱即用** — 装完依赖直接用，无需复杂配置\n- 📦 **支持压缩包** — 不用手动解压，直接列出并读取内部文件\n- 🔍 **模糊匹配** — 大小写不敏感匹配文件名，找不到精确匹配时自动尝试\n- 🎨 **多种输出格式** — 人类可读文本 / JSON 程序接口都支持\n- 🧩 **完整支持所有常用格式** — 办公文档+压缩包全覆盖\n\n## 📝 使用场景\n\n- AI 分析各种办公文档\n- 批量读取压缩包内的文档内容\n- 快速查看附件内容\n- 数据提取和预处理\n\n## 作者\n\nCreated by xiaoya Liu with OpenClaw\n\nFile v1.0.0:_meta.json\n\n{\n  \"ownerId\": \"kn7bpabfda1pxgynfp9w46kjjd84fe8z\",\n  \"slug\": \"document-reader\",\n  \"version\": \"1.0.0\",\n  \"publishedAt\": 1775637768169\n}\n\nFile v1.0.0:skill-card.md\n\n## Description:\n\nDocument Reader extracts text from PDF, DOCX, XLSX, PPTX, RTF, ODT, HTML, plain text, and selected files inside common archives for agent analysis.\n\nThis skill is ready for commercial/non-commercial use.\n\n## Publisher:\n\n[xiaoyaliu00](https://clawhub.ai/user/xiaoyaliu00)\n\n### License/Terms of Use:\n\nMIT-0\n\n## Use Case:\n\nDevelopers, engineers, and agent users use this skill to read common office documents, text files, and files stored inside archives so the extracted content can be analyzed or processed by an agent.\n\n### Deployment Geography for Use:\n\nGlobal\n\n## Known Risks and Mitigations:\n\nRisk: Archive handling and document extraction can expose the agent environment to untrusted or oversized files.\n\nMitigation: Use the skill only on files the user intends the agent to read, avoid untrusted oversized archives, and review extracted content before relying on it.\n\nRisk: The documented dependency installation command can install packages system-wide.\n\nMitigation: Install dependencies in an isolated virtual environment or container before running the document reader.\n\n## Reference(s):\n\n\n## Skill Output:\n\n**Output Type(s):** [text, markdown, json, shell commands, guidance]\n\n**Output Format:** [Plain text or JSON emitted by a command-line document reader, with Markdown documentation and shell command examples.]\n\n**Output Parameters:** [1D]\n\n**Other Properties Related to Output:** [Can list archive contents or read a specified file from an archive; long text output may be truncated in human-readable mode.]\n\n## Skill Version(s):\n\n1.0.0 (source: frontmatter and server-resolved release metadata)\n\n## Ethical Considerations:\n\nUsers should evaluate whether this skill is appropriate for their environment, review any generated or modified files before relying on them, and apply their organization's safety, security, and compliance requirements before deployment.","readmeExcerpt":"Skill: document-reader Owner: xiaoyaliu00 Summary: 通用文档读取工具，支持 PDF/DOCX/XLSX/PPTX/RTF/ODT 等多种文档格式，也支持 ZIP/TAR.GZ/RAR/7Z 等主流压缩包内文档直接读取 Tags: latest:1.0.0 Version history: v1.0.0 | 2026-04-08T08:42:48.169Z | user Initial release: support PDF/DOCX/XLSX/PPTX/RTF/ODT documents and ZIP/TAR/RAR/7Z compressed files Archive index: Archive v1.0.0: 4 files, 6842 bytes Files: scripts/document_reader.py (17331b), skill-card.md (1","codeSnippets":[],"executableExamples":[{"language":"bash","snippet":"pip install textract python-docx openpyxl python-pptx rarfile py7zr --break-system-packages"},{"language":"bash","snippet":"apt-get install -y poppler-utils antiword unrtf tidy libxml2-dev libxslt1-dev"},{"language":"bash","snippet":"# 读取 PDF 文件\npython {baseDir}/scripts/document_reader.py --file /path/to/document.pdf\n\n# 读取 Word 文档\npython {baseDir}/scripts/document_reader.py --file /path/to/report.docx\n\n# 读取 Excel 文件（输出带 Sheet 分隔的表格文本）\npython {baseDir}/scripts/document_reader.py --file /path/to/data.xlsx\n\n# JSON 格式输出（方便程序处理）\npython {baseDir}/scripts/document_reader.py --file /path/to/data.xlsx --format json"},{"language":"bash","snippet":"# 列出 ZIP 包内容\npython {baseDir}/scripts/document_reader.py --list /path/to/archive.zip\n\n# 列出 RAR 包内容\npython {baseDir}/scripts/document_reader.py --list /path/to/archive.rar\n\n# 列出 7z 包内容\npython {baseDir}/scripts/document_reader.py --list /path/to/archive.7z"},{"language":"bash","snippet":"# 读取 ZIP 包内的 Word 文档\npython {baseDir}/scripts/document_reader.py --file /path/to/archive.zip --inner-path document.docx\n\n# 读取 7z 包内的 PDF\npython {baseDir}/scripts/document_reader.py --file /path/to/archive.7z --inner-path report.pdf"},{"language":"text","snippet":"=== report.pdf ===\n\n# 项目进度报告\n\n## 本周完成\n\n1. 完成了前端界面开发\n2. 修复了三个 Bug\n3. 编写了接口文档\n\n..."}],"parameters":null,"dependencies":[],"permissions":[],"extractedFiles":[{"path":"SKILL.md","content":"---\nname: document-reader\ndescription: \"通用文档读取工具，支持 PDF/DOCX/XLSX/PPTX/RTF/ODT 等多种文档格式，也支持 ZIP/TAR.GZ/RAR/7Z 等主流压缩包内文档直接读取\"\nversion: \"1.0.0\"\n---\n\n# Document Reader - 通用文档读取技能\n\n读取各种格式文档的内容，直接输出文本供 AI 分析。支持压缩包内文档直接读取，无需手动解压。\n\n## 📦 支持格式\n\n### 文档\n| 格式 | 扩展名 | 说明 |\n|------|--------|------|\n| PDF | `.pdf` | 支持文本提取，依赖 poppler-utils |\n| Microsoft Word | `.docx` | 完整提取所有段落文本 |\n| Microsoft Excel | `.xlsx` | 按 Sheet 输出，每个 Sheet 输出为表格文本 |\n| Microsoft PowerPoint | `.pptx` | 按 Slide 分块输出 |\n| Rich Text Format | `.rtf` |  |\n| OpenDocument Text | `.odt` |  |\n| HTML | `.html`/`.htm` | 提取正文文本 |\n| 纯文本 | `.txt`/`.md`/`.json`/`.xml`/`.py`/`.js` 等 | 直接读取 |\n\n### 压缩包\n| 格式 | 扩展名 | 功能 |\n|------|--------|------|\n| ZIP | `.zip` | 列出文件 ➜ 读取指定文档 |\n| TAR | `.tar`/`.tar.gz`/`.tgz`/`.tar.bz2` | 列出文件 ➜ 读取指定文档 |\n| RAR | `.rar` | 列出文件 ➜ 读取指定文档 |\n| 7-Zip | `.7z` | 列出文件 ➜ 读取指定文档 |\n\n## 🚀 快速开始\n\n### 依赖安装\n\n**Python 包：**\n```bash\npip install textract python-docx openpyxl python-pptx rarfile py7zr --break-system-packages\n```\n\n**系统依赖（Ubuntu/Debian）：**\n```bash\napt-get install -y poppler-utils antiword unrtf tidy libxml2-dev libxslt1-dev\n```\n\n## 💡 使用示例\n\n### 1. 直接读取本地文档\n\n```bash\n# 读取 PDF 文件\npython {baseDir}/scripts/document_reader.py --file /path/to/document.pdf\n\n# 读取 Word 文档\npython {baseDir}/scripts/document_reader.py --file /path/to/report.docx\n\n# 读取 Excel 文件（输出带 Sheet 分隔的表格文本）\npython {baseDir}/scripts/document_reader.py --file /path/to/data.xlsx\n\n# JSON 格式输出（方便程序处理）\npython {baseDir}/scripts/document_reader.py --file /path/to/data.xlsx --format json\n```\n\n### 2. 处理压缩包\n\n**先列出压缩包里有哪些文件：**\n```bash\n# 列出 ZIP 包内容\npython {baseDir}/scripts/document_reader.py --list /path/to/archive.zip\n\n# 列出 RAR 包内容\npython {baseDir}/scripts/document_reader.py --list /path/to/archive.rar\n\n# 列出 7z 包内容\npython {baseDir}/scripts/document_reader.py --list /path/to/archive.7z\n```\n\n**读取压缩包里的指定文档：**\n```bash\n# 读取 ZIP 包内的 Word 文档\npython {baseDir}/scripts/document_reader.py --file /path/to/archive.zip --inner-path document.docx\n\n# 读取 7z 包内的 PDF\npython {baseDir}/scripts/document_reader.py --file /path/to/archive.7z --inner-path report.pdf\n```\n\n### 输出示例\n\n**读取文档：**\n```\n=== report.pdf ===\n\n# 项目进度报告\n\n## 本周完成\n\n1. 完成了前端界面开发\n2. 修复了三个 Bug\n3. 编写了接口文档\n\n...\n```\n\n**列出压缩包：**\n```\nArchive: data.zip\nFound 3 file(s):\n\n  readme.txt\n  docs/report.pdf\n  data/sheet.xlsx\n```\n\n## ✨ 特性\n\n- 🎯 **开箱即用** — 装完依赖直接用，无需复杂配置\n- 📦 **支持压缩包** — 不用手动解压，直接列出并读取内部文件\n- 🔍 **模糊匹配** — 大小写不敏感匹配文件名，找不到精确匹配时自动尝试\n- 🎨 **多种输出格式** — 人类可读文本 / JSON 程序接口都支持\n- 🧩 **完整支持所有常用格式** — 办公文档+压缩包全覆盖\n\n## 📝 使用场景\n\n- AI 分析各种办公文档\n- 批量读取压缩包内的文档内容\n- 快速查看附件内容\n- 数据提取和预处理\n\n## 作者\n\nCreated by xiaoya Liu with OpenClaw"},{"path":"_meta.json","content":"{\n  \"ownerId\": \"kn7bpabfda1pxgynfp9w46kjjd84fe8z\",\n  \"slug\": \"document-reader\",\n  \"version\": \"1.0.0\",\n  \"publishedAt\": 1775637768169\n}"},{"path":"skill-card.md","content":"## Description:\n\nDocument Reader extracts text from PDF, DOCX, XLSX, PPTX, RTF, ODT, HTML, plain text, and selected files inside common archives for agent analysis.\n\nThis skill is ready for commercial/non-commercial use.\n\n## Publisher:\n\n[xiaoyaliu00](https://clawhub.ai/user/xiaoyaliu00)\n\n### License/Terms of Use:\n\nMIT-0\n\n## Use Case:\n\nDevelopers, engineers, and agent users use this skill to read common office documents, text files, and files stored inside archives so the extracted content can be analyzed or processed by an agent.\n\n### Deployment Geography for Use:\n\nGlobal\n\n## Known Risks and Mitigations:\n\nRisk: Archive handling and document extraction can expose the agent environment to untrusted or oversized files.\n\nMitigation: Use the skill only on files the user intends the agent to read, avoid untrusted oversized archives, and review extracted content before relying on it.\n\nRisk: The documented dependency installation command can install packages system-wide.\n\nMitigation: Install dependencies in an isolated virtual environment or container before running the document reader.\n\n## Reference(s):\n\n\n## Skill Output:\n\n**Output Type(s):** [text, markdown, json, shell commands, guidance]\n\n**Output Format:** [Plain text or JSON emitted by a command-line document reader, with Markdown documentation and shell command examples.]\n\n**Output Parameters:** [1D]\n\n**Other Properties Related to Output:** [Can list archive contents or read a specified file from an archive; long text output may be truncated in human-readable mode.]\n\n## Skill Version(s):\n\n1.0.0 (source: frontmatter and server-resolved release metadata)\n\n## Ethical Considerations:\n\nUsers should evaluate whether this skill is appropriate for their environment, review any generated or modified files before relying on them, and apply their organization's safety, security, and compliance requirements before deployment."}],"languages":[],"docsSourceLabel":"CLAWHUB","editorialOverview":"通用文档读取工具，支持 PDF/DOCX/XLSX/PPTX/RTF/ODT 等多种文档格式，也支持 ZIP/TAR.GZ/RAR/7Z 等主流压缩包内文档直接读取 Skill: document-reader Owner: xiaoyaliu00 Summary: 通用文档读取工具，支持 PDF/DOCX/XLSX/PPTX/RTF/ODT 等多种文档格式，也支持 ZIP/TAR.GZ/RAR/7Z 等主流压缩包内文档直接读取 Tags: latest:1.0.0 Version history: v1.0.0 | 2026-04-08T08:42:48.169Z | user Initial release: support PDF/DOCX/XLSX/PPTX/RTF/ODT documents and ZIP/TAR/RAR/7Z compressed files Archive index: Archive v1.0.0: 4 files, 6842 bytes Files: scripts/document_reader.py (17331b), skill-card.md (1","editorialQuality":{"score":100,"threshold":65,"status":"ready","wordCount":730,"uniquenessScore":57,"reasons":[]}},"media":{"evidence":{"source":"no-media","verified":false,"confidence":"low","updatedAt":"2026-10-09T05:56:35.991Z","emptyReason":"No screenshots, media assets, or demo links are available."},"primaryImageUrl":null,"mediaAssetCount":0,"assets":[],"demoUrl":null},"ownerResources":{"evidence":{"source":"unclaimed","verified":false,"confidence":"low","updatedAt":"2026-10-09T05:56:35.991Z","emptyReason":"This page has not been claimed by the agent owner."},"hasCustomPage":false,"customPageUpdatedAt":null,"customLinks":[],"structuredLinks":{"docsUrl":null,"demoUrl":null,"supportUrl":null,"pricingUrl":null,"statusUrl":null},"customPage":null},"relatedAgents":{"evidence":{"source":"protocol-neighbors","verified":false,"confidence":"medium","updatedAt":"2026-10-09T06:09:42.373Z","emptyReason":null},"items":[{"id":"b917f68a-ebff-438e-84f8-3f4b2494c0bc","entityType":"agent","canonicalPath":"/agent/activepieces-activepieces","slug":"activepieces-activepieces","name":"activepieces","description":"AI Agents & MCPs & AI Workflow Automation • (~400 MCP servers for AI agents) • AI Automation / AI Agent with MCPs • AI Workflows & AI Agents • MCPs for AI Agents","url":"https://github.com/activepieces/activepieces","homepage":"https://www.activepieces.com","source":"GITHUB_REPOS","protocols":["OPENCLAW"],"capabilities":[],"safetyScore":100,"overallRank":70,"updatedAt":"2026-04-15T02:22:12.426Z","createdAt":"2026-02-25T03:38:12.412Z","downloads":null},{"id":"5cb26759-3a39-483f-94cf-276a98c13bb8","entityType":"agent","canonicalPath":"/agent/cherryhq-cherry-studio","slug":"cherryhq-cherry-studio","name":"cherry-studio","description":"AI productivity studio with smart chat, autonomous agents, and 300+ assistants. Unified access to frontier LLMs","url":"https://github.com/CherryHQ/cherry-studio","homepage":"https://cherry-ai.com","source":"GITHUB_REPOS","protocols":["MCP","OPENCLAW"],"capabilities":[],"safetyScore":100,"overallRank":70,"updatedAt":"2026-04-11T14:38:40.986Z","createdAt":"2026-02-25T03:38:19.379Z","downloads":null},{"id":"8ebccd8e-3863-4187-8355-c3f14e1f9edf","entityType":"agent","canonicalPath":"/agent/iofficeai-aionui","slug":"iofficeai-aionui","name":"AionUi","description":"Free, local, open-source 24/7 Cowork app and OpenClaw for Gemini CLI, Claude Code, Codex, OpenCode, Qwen Code, Goose CLI, Auggie, and more | 🌟 Star if you like it!","url":"https://github.com/iOfficeAI/AionUi","homepage":"https://www.aionui.com","source":"GITHUB_REPOS","protocols":["MCP","OPENCLAW"],"capabilities":[],"safetyScore":100,"overallRank":70,"updatedAt":"2026-04-10T18:48:31.762Z","createdAt":"2026-02-25T03:38:16.584Z","downloads":null},{"id":"6f6582d0-5d76-4f0f-b81d-86520247950b","entityType":"agent","canonicalPath":"/agent/copilotkit-copilotkit","slug":"copilotkit-copilotkit","name":"CopilotKit","description":"The Frontend for Agents & Generative UI. React + Angular","url":"https://github.com/CopilotKit/CopilotKit","homepage":"https://docs.copilotkit.ai","source":"GITHUB_REPOS","protocols":["OPENCLAW"],"capabilities":[],"safetyScore":100,"overallRank":70,"updatedAt":"2026-03-25T09:50:57.846Z","createdAt":"2026-02-25T03:39:14.617Z","downloads":null}],"links":{"hub":"/agent","source":"/agent/source/clawhub","protocols":[{"label":"OpenClaw","href":"/agent/protocol/openclew"}]}}}