{"id":"c123b053-ae0a-4047-8902-f158366bf170","entityType":"agent","slug":"clawhub-dinstein-tech-news-digest","name":"Tech News Digest","canonicalUrl":"https://www.xpersona.co/agent/clawhub-dinstein-tech-news-digest","canonicalPath":"/agent/clawhub-dinstein-tech-news-digest","generatedAt":"2026-10-09T14:02:08.759Z","source":"CLAWHUB","claimStatus":"UNCLAIMED","verificationTier":"NONE","summary":{"evidence":{"source":"editorial-content","verified":true,"confidence":"high","updatedAt":"2026-04-15T00:45:39.800Z","emptyReason":null},"description":"Generate tech news digests with unified source model, quality scoring, and multi-format output. Five-layer data collection from RSS feeds, Twitter/X KOLs, Gi... Skill: Tech News Digest Owner: dinstein Summary: Generate tech news digests with unified source model, quality scoring, and multi-format output. Five-layer data collection from RSS feeds, Twitter/X KOLs, Gi... Tags: latest:3.11.0 Version history: v3.11.0 | 2026-02-28T16:22:08.196Z | user Tavily backend, quality scores, domain limit fix, tests, CI v3.10.3 | 2026-02-27T15:46:08.147Z | user Docs alignment + config namin","descriptionLabel":"Technical summary","evidenceSummary":"Capability contract not published. No trust telemetry is available yet. 2K downloads reported by the source. Last updated 4/15/2026.","installCommand":"clawhub skill install kn74589cx1nbhnc3x0f3nwre39814699:tech-news-digest","sourceUrl":"https://clawhub.ai/dinstein/tech-news-digest","homepage":"https://clawhub.ai/dinstein/tech-news-digest","primaryLinks":[{"label":"View on ClawHub","url":"https://clawhub.ai/dinstein/tech-news-digest","kind":"source"}],"safetyScore":84,"overallRank":62,"popularityScore":66,"trustScore":null,"claimedByName":null,"isOwner":false,"seoDescription":"Generate tech news digests with unified source model, quality scoring, and multi-format output. Five-layer data collection from RSS feeds, Twitter/X KOLs, Gi..."},"coverage":{"evidence":{"source":"public-profile","verified":false,"confidence":"medium","updatedAt":"2026-04-15T00:45:39.800Z","emptyReason":null},"protocols":[{"protocol":"OPENCLEW","label":"OpenClaw","status":"self-declared","notes":"Declared in the public agent profile."}],"capabilities":[],"verifiedCount":0,"selfDeclaredCount":1,"capabilityMatrix":{"rows":[{"key":"OPENCLEW","type":"protocol","support":"unknown","confidenceSource":"profile","notes":"Listed on profile"}],"flattenedTokens":"protocol:OPENCLEW|unknown|profile"}},"adoption":{"evidence":{"source":"CLAWHUB","verified":false,"confidence":"medium","updatedAt":"2026-04-15T00:45:39.800Z","emptyReason":null},"stars":null,"forks":null,"downloads":2039,"packageName":null,"latestVersion":"3.11.0","tractionLabel":"2K downloads"},"release":{"evidence":{"source":"CLAWHUB","verified":false,"confidence":"medium","updatedAt":"2026-02-28T19:23:04.645Z","emptyReason":null},"lastUpdatedAt":"2026-04-15T00:45:39.800Z","lastCrawledAt":"2026-02-28T19:23:04.645Z","lastIndexedAt":null,"nextCrawlAt":"2026-03-01T19:23:04.645Z","lastVerifiedAt":null,"highlights":[{"version":"3.11.0","createdAt":"2026-02-28T16:22:08.196Z","changelog":"Tavily backend, quality scores, domain limit fix, tests, CI","fileCount":37,"zipByteSize":113229},{"version":"3.10.3","createdAt":"2026-02-27T15:46:08.147Z","changelog":"Docs alignment + config naming refactor","fileCount":29,"zipByteSize":96907},{"version":"3.10.2","createdAt":"2026-02-27T12:01:25.639Z","changelog":"Fix domain limits for multi-author platforms + Brave multi-key rotation","fileCount":null,"zipByteSize":null},{"version":"3.6.2","createdAt":"2026-02-22T06:23:22.415Z","changelog":"Add 3 GitHub sources: cloudflare/moltworker, sipeed/picoclaw, HKUDS/nanobot","fileCount":null,"zipByteSize":null},{"version":"3.6.1","createdAt":"2026-02-22T06:17:37.818Z","changelog":"v3.6.1: prompt review pass","fileCount":null,"zipByteSize":null},{"version":"3.6.0","createdAt":"2026-02-21T06:57:08.561Z","changelog":"v3.6.0","fileCount":null,"zipByteSize":null},{"version":"3.5.1","createdAt":"2026-02-21T06:28:06.959Z","changelog":"v3.5.1","fileCount":null,"zipByteSize":null},{"version":"3.5.0","createdAt":"2026-02-18T14:47:10.890Z","changelog":"P1 fixes: per-run cache isolation, per-topic domain limits, fix 6 Twitter handles, zsh compat","fileCount":null,"zipByteSize":null}]},"execution":{"evidence":{"source":"CLAWHUB","verified":false,"confidence":"low","updatedAt":null,"emptyReason":"No published capability contract is available yet."},"installCommand":"clawhub skill install kn74589cx1nbhnc3x0f3nwre39814699:tech-news-digest","setupComplexity":"low","setupSteps":["Setup complexity is classified as HIGH. You must provision dedicated cloud infrastructure or an isolated VM. Do not run this directly on your local workstation.","Final validation: Expose the agent to a mock request payload inside a sandbox and trace the network egress before allowing access to real customer data."],"contract":{"contractStatus":"missing","authModes":[],"requires":[],"forbidden":[],"supportsMcp":false,"supportsA2a":false,"supportsStreaming":false,"inputSchemaRef":null,"outputSchemaRef":null,"dataRegion":null,"contractUpdatedAt":null,"sourceUpdatedAt":null,"freshnessSeconds":null},"invocationGuide":{"preferredApi":{"snapshotUrl":"https://www.xpersona.co/api/v1/agents/clawhub-dinstein-tech-news-digest/snapshot","contractUrl":"https://www.xpersona.co/api/v1/agents/clawhub-dinstein-tech-news-digest/contract","trustUrl":"https://www.xpersona.co/api/v1/agents/clawhub-dinstein-tech-news-digest/trust"},"curlExamples":["curl -s \"https://www.xpersona.co/api/v1/agents/clawhub-dinstein-tech-news-digest/snapshot\"","curl -s \"https://www.xpersona.co/api/v1/agents/clawhub-dinstein-tech-news-digest/contract\"","curl -s \"https://www.xpersona.co/api/v1/agents/clawhub-dinstein-tech-news-digest/trust\""],"jsonRequestTemplate":{"query":"summarize this repo","constraints":{"maxLatencyMs":2000,"protocolPreference":["OPENCLEW"]}},"jsonResponseTemplate":{"ok":true,"result":{"summary":"...","confidence":0.9},"meta":{"source":"CLAWHUB","generatedAt":"2026-10-09T14:02:08.754Z"}},"retryPolicy":{"maxAttempts":3,"backoffMs":[500,1500,3500],"retryableConditions":["HTTP_429","HTTP_503","NETWORK_TIMEOUT"]}},"endpoints":{"dossierUrl":"https://www.xpersona.co/api/v1/agents/clawhub-dinstein-tech-news-digest/dossier","snapshotUrl":"https://www.xpersona.co/api/v1/agents/clawhub-dinstein-tech-news-digest/snapshot","contractUrl":"https://www.xpersona.co/api/v1/agents/clawhub-dinstein-tech-news-digest/contract","trustUrl":"https://www.xpersona.co/api/v1/agents/clawhub-dinstein-tech-news-digest/trust"}},"reliability":{"evidence":{"source":"runtime-metrics","verified":false,"confidence":"low","updatedAt":null,"emptyReason":"No trust, reliability, or runtime telemetry is available."},"trust":{"status":"unavailable","handshakeStatus":"UNKNOWN","verificationFreshnessHours":null,"reputationScore":null,"p95LatencyMs":null,"successRate30d":null,"fallbackRate":null,"attempts30d":null,"trustUpdatedAt":null,"trustConfidence":"unknown","sourceUpdatedAt":null,"freshnessSeconds":null},"decisionGuardrails":{"doNotUseIf":["Contract metadata is missing or unavailable for deterministic execution."],"safeUseWhen":[],"riskFlags":["missing_or_unavailable_contract","trust_data_unavailable","schema_references_missing"],"operationalConfidence":"low"},"executionMetrics":{"observedLatencyMsP50":null,"observedLatencyMsP95":null,"estimatedCostUsd":null,"uptime30d":null,"rateLimitRpm":null,"rateLimitBurst":null,"lastVerifiedAt":null,"verificationSource":null},"runtimeMetrics":{"successRate":null,"avgLatencyMs":null,"avgCostUsd":null,"hallucinationRate":null,"retryRate":null,"disputeRate":null,"p50Latency":null,"p95Latency":null,"lastUpdated":null}},"benchmarks":{"evidence":{"source":"no-benchmark-data","verified":false,"confidence":"low","updatedAt":null,"emptyReason":"No benchmark suites or observed failure patterns are available."},"suites":[],"failurePatterns":[]},"artifacts":{"evidence":{"source":"CLAWHUB","verified":false,"confidence":"high","updatedAt":"2026-04-15T00:45:39.800Z","emptyReason":null},"readme":"Skill: Tech News Digest\n\nOwner: dinstein\n\nSummary: Generate tech news digests with unified source model, quality scoring, and multi-format output. Five-layer data collection from RSS feeds, Twitter/X KOLs, Gi...\n\nTags: latest:3.11.0\n\nVersion history:\n\nv3.11.0 | 2026-02-28T16:22:08.196Z | user\n\nTavily backend, quality scores, domain limit fix, tests, CI\n\nv3.10.3 | 2026-02-27T15:46:08.147Z | user\n\nDocs alignment + config naming refactor\n\nv3.10.2 | 2026-02-27T12:01:25.639Z | user\n\nFix domain limits for multi-author platforms + Brave multi-key rotation\n\nv3.6.2 | 2026-02-22T06:23:22.415Z | user\n\nAdd 3 GitHub sources: cloudflare/moltworker, sipeed/picoclaw, HKUDS/nanobot\n\nv3.6.1 | 2026-02-22T06:17:37.818Z | user\n\nv3.6.1: prompt review pass\n\nv3.6.0 | 2026-02-21T06:57:08.561Z | user\n\nv3.6.0\n\nv3.5.1 | 2026-02-21T06:28:06.959Z | user\n\nv3.5.1\n\nv3.5.0 | 2026-02-18T14:47:10.890Z | user\n\nP1 fixes: per-run cache isolation, per-topic domain limits, fix 6 Twitter handles, zsh compat\n\nv3.4.11 | 2026-02-18T13:48:07.508Z | user\n\nFix 6 Twitter handles, zsh compat for test-pipeline, github token debug logs\n\nv3.4.10 | 2026-02-18T00:46:12.449Z | user\n\nRepublish\n\nv3.4.9 | 2026-02-17T15:37:13.956Z | auto\n\nVersion 3.4.9\n\n- Added summarize-merged.py script for output summarization.\n- Updated SKILL.md to include \"openssl\" as an optional dependency.\n- Documentation improvements and updates to changelog and reference prompts.\n\nv3.4.7 | 2026-02-17T14:56:10.812Z | auto\n\nVersion 3.4.7\n\n- Updated `fetch-github.py` script.\n- Documentation improvements in SKILL.md, including clarified template and output details. \n- Version bump and changelog maintenance.\n\nv3.4.6 | 2026-02-17T14:08:50.097Z | auto\n\n- Improved security note for GH_APP_TOKEN_SCRIPT: now warns users that only trusted scripts should be referenced, as the script will be executed via subprocess.\n- Updated all documentation and schema references to clarify workspace and archive directory naming (now references 'tech-news-digest' instead of 'tech-digest').\n- Corrected environment variable documentation for generating GitHub App tokens.\n- Minor documentation and schema consistency updates across config, template, and reference files for clarity and safety.\n\nv3.4.5 | 2026-02-17T10:47:05.936Z | auto\n\n- Improved reliability and internal consistency across all data collection scripts.\n- Updated documentation in SKILL.md, README.md, and CHANGELOG.md for recent enhancements.\n- Refactored source configuration loader and validation for clearer error handling.\n- Enhanced logging and status reporting in fetch and merge scripts.\n- Minor bug fixes and code clean-up.\n\nv3.4.4 | 2026-02-17T10:35:49.853Z | auto\n\n## tech-news-digest 3.4.4\n\n- Updated version to 3.4.4 in documentation.\n- No changes to code or behavior; documentation/version consistency update only.\n\nv3.4.3 | 2026-02-17T09:41:09.209Z | auto\n\nVersion 3.4.3\n\n- Added \"gh\" to the list of optional CLI tools in the requirements.\n- No functional or pipeline changes; documentation and metadata updated.\n\nv3.4.2 | 2026-02-17T09:39:10.447Z | auto\n\n- Added new environment variables for GitHub App authentication: GH_APP_ID, GH_APP_INSTALL_ID, GH_APP_KEY_FILE, and GH_APP_TOKEN_SCRIPT.\n- Updated documentation to reflect additional GitHub App token generation options.\n- Improved GitHub token generation logic in scripts/fetch-github.py to support new environment variables for automatic installation token creation.\n\nv3.4.1 | 2026-02-17T07:00:20.427Z | auto\n\n- Improved merge-sources.py script to handle additional merging logic for topic and source overrides.\n- Updated default sources.json with new tech sources and refined topic mappings.\n- Enhanced Discord, email, and Telegram template files for clearer digest formatting.\n- Minor documentation fixes and clarifications in SKILL.md and prompts.\n- Improved configuration validation and documentation for workspace overrides.\n\nv3.4.0 | 2026-02-17T01:44:28.875Z | auto\n\ntech-news-digest 3.4.0 introduces a unified pipeline script and improved programmatic control.\n\n- Added scripts/run-pipeline.py: a new single-entry script running all collectors in parallel and merging results.\n- Updated all fetch scripts to accept unified --defaults and --config arguments.\n- Auth improvements: GitHub token now auto-generates from a GitHub App if $GITHUB_TOKEN is unset.\n- Enhanced configuration flexibility and environment variable descriptions in documentation.\n- README and SKILL.md revised to highlight the unified workflow and clarify recommended usage.\n\nv3.3.2 | 2026-02-16T16:00:33.444Z | user\n\nDeclare tools (python3, gog) and file access paths in metadata; addresses VT undeclared tools audit\n\nv3.3.1 | 2026-02-16T15:53:58.637Z | user\n\nRemove anthropic-rss third-party mirror (supply chain risk); update source count to 131\n\nv3.2.0 | 2026-02-16T03:30:14.448Z | auto\n\nVersion 3.2.0\n\n- Updated and improved output templates for Discord, email, Markdown, and Telegram.\n- Enhanced reference prompt and documentation in reference files.\n- Minor documentation and changelog updates for clarity and accuracy.\n- No changes to core pipeline logic or configuration schema.\n\nv3.1.0 | 2026-02-16T03:23:00.182Z | auto\n\n- Added detailed inline documentation and usage examples for all supported output templates (Discord, email, Telegram) and digest prompts.\n- Improved template files for Discord, email, and Telegram to clarify formatting and message structure.\n- Updated references and prompt documentation to streamline multi-format output customization.\n- Expanded SKILL.md with clearer template usage instructions and examples.\n- General content and documentation enhancements for usability and clarity.\n\nv3.0.0 | 2026-02-16T03:02:23.033Z | auto\n\nMajor update: Reddit support and expanded source coverage.\n\n- Added Reddit as a new data layer with `fetch-reddit.py`, enabling five-layer news aggregation.\n- Increased default sources to 132, including 13 major Reddit communities.\n- Updated config and pipeline scripts to integrate Reddit alongside RSS, Twitter/X, GitHub, and web search.\n- Enhanced SKILL.md, documentation, and example configs to reflect the new Reddit integration.\n- Improved source model and topic tagging for broader and more customizable tech news coverage.\n\nv2.8.1 | 2026-02-16T01:55:39.092Z | auto\n\n- Refreshed digest templates for Discord, email, and Telegram with improved formatting and metadata.\n- Updated default source list and topic definitions for better coverage and tag accuracy.\n- Revised documentation in SKILL.md for clarity and up-to-date usage instructions.\n- Enhanced prompt guidance in references/digest-prompt.md.\n- Minor improvements to default config (sources.json), template usability, and config instructions.\n\nv2.7.0 | 2026-02-16T00:36:45.486Z | auto\n\ntech-news-digest v2.7.0\n\n- Updated prompt and template files for improved digest formatting and clarity.\n- Refined Discord, email, and Telegram output templates for more consistent multi-platform presentation.\n- Enhanced reference documentation for prompt design and template usage.\n- Minor content and metadata adjustments in skill documentation.\n\nv2.6.1 | 2026-02-16T00:09:56.542Z | auto\n\n- Improved Twitter/X KOL collection: Now skips over empty user handles for better data integrity.\n- Minor update to documentation in SKILL.md.\n- Patch version bump to 2.6.1.\n\nv2.6.0 | 2026-02-16T00:04:39.583Z | auto\n\n- Adds a new config option to set both defaults and workspace paths in the config validator (`validate-config.py`).\n- Updates `validate-config.py` argument interface to `--defaults` / `--config` for better clarity (was formerly `--config-dir`).\n- Updates documentation and usage examples in SKILL.md for consistent interface references.\n- Enhances configuration validation flexibility and accuracy.\n- Minor documentation cleanup and consistency fixes across references and scripts.\n\nv2.5.0 | 2026-02-15T16:46:14.680Z | user\n\nFix twitter reply filter, scoring consistency, template versions, article count, pipeline resume, source health monitoring, e2e test, archive cleanup, rate limiting, web scoring, dead code removal\n\nv2.4.1 | 2026-02-15T16:02:18.260Z | user\n\nRemove version from SKILL.md title, move changelog to CHANGELOG.md\n\nv2.4.0 | 2026-02-15T15:55:35.351Z | user\n\nPerformance: batch Twitter lookup, smart dedup, conditional fetch for RSS/GitHub\n\nv2.2.1 | 2026-02-15T14:36:27.427Z | user\n\nSecurity hardening: shell injection prevention for email delivery, added Security Notes section documenting execution model, network access, file access scope\n\nv2.2.0 | 2026-02-15T12:17:30.123Z | user\n\nRenamed from tech-digest. 109 sources, 4-layer pipeline (RSS/Twitter/Web/GitHub), quality scoring, multi-template output, cron prompt pattern.\n\nArchive index:\n\nArchive v3.11.0: 37 files, 113229 bytes\n\nFiles: CHANGELOG.md (17649b), config/defaults/sources.json (38686b), config/defaults/topics.json (2951b), config/schema.json (4561b), CONTRIBUTING.md (2283b), README_CN.md (4592b), README.md (4583b), references/digest-prompt.md (5933b), references/templates/discord.md (3011b), references/templates/email.md (5805b), references/templates/pdf.md (2119b), requirements.txt (606b), scripts/config_loader.py (8490b), scripts/fetch-github.py (20476b), scripts/fetch-reddit.py (12211b), scripts/fetch-rss.py (19352b), scripts/fetch-twitter.py (28449b), scripts/fetch-web.py (26778b), scripts/generate-pdf.py (9751b), scripts/merge-sources.py (24247b), scripts/run-pipeline.py (9960b), scripts/sanitize-html.py (7357b), scripts/send-email.py (5199b), scripts/source-health.py (4931b), scripts/summarize-merged.py (3526b), scripts/test-pipeline.sh (10975b), scripts/validate-config.py (9659b), SKILL.md (21091b), tests/fixtures/github.json (3488b), tests/fixtures/merged.json (13075b), tests/fixtures/reddit.json (6097b), tests/fixtures/rss.json (3374b), tests/fixtures/twitter.json (5039b), tests/fixtures/web.json (4138b), tests/test_config.py (4574b), tests/test_merge.py (9779b), _meta.json (136b)\n\nFile v3.11.0:SKILL.md\n\n---\nname: tech-news-digest\ndescription: Generate tech news digests with unified source model, quality scoring, and multi-format output. Five-layer data collection from RSS feeds, Twitter/X KOLs, GitHub releases, Reddit, and web search. Pipeline-based scripts with retry mechanisms and deduplication. Supports Discord, email, and markdown templates.\nversion: \"3.11.0\"\nhomepage: https://github.com/draco-agent/tech-news-digest\nsource: https://github.com/draco-agent/tech-news-digest\nmetadata:\n  openclaw:\n    requires:\n      bins: [\"python3\"]\n    optionalBins: [\"mail\", \"msmtp\", \"gog\", \"gh\", \"openssl\", \"weasyprint\"]\nenv:\n  - name: TWITTER_API_BACKEND\n    required: false\n    description: \"Twitter API backend: 'official', 'twitterapiio', or 'auto' (default: auto)\"\n  - name: X_BEARER_TOKEN\n    required: false\n    description: Twitter/X API bearer token for KOL monitoring (official backend)\n  - name: TWITTERAPI_IO_KEY\n    required: false\n    description: twitterapi.io API key for KOL monitoring (twitterapiio backend)\n  - name: TAVILY_API_KEY\n    required: false\n    description: Tavily Search API key (alternative to Brave)\n  - name: WEB_SEARCH_BACKEND\n    required: false\n    description: \"Web search backend: auto (default), brave, or tavily\"\n  - name: BRAVE_API_KEYS\n    required: false\n    description: Brave Search API keys (comma-separated for rotation)\n  - name: BRAVE_API_KEY\n    required: false\n    description: Brave Search API key (single key fallback)\n  - name: GITHUB_TOKEN\n    required: false\n    description: GitHub token for higher API rate limits (auto-generated from GitHub App if not set)\n  - name: GH_APP_ID\n    required: false\n    description: GitHub App ID for automatic installation token generation\n  - name: GH_APP_INSTALL_ID\n    required: false\n    description: GitHub App Installation ID for automatic token generation\n  - name: GH_APP_KEY_FILE\n    required: false\n    description: Path to GitHub App private key PEM file\ntools:\n  - python3: Required. Runs data collection and merge scripts.\n  - mail: Optional. msmtp-based mail command for email delivery (preferred).\n  - gog: Optional. Gmail CLI for email delivery (fallback if mail not available).\nfiles:\n  read:\n    - config/defaults/: Default source and topic configurations\n    - references/: Prompt templates and output templates\n    - scripts/: Python pipeline scripts\n    - <workspace>/archive/tech-news-digest/: Previous digests for dedup\n  write:\n    - /tmp/td-*.json: Temporary pipeline intermediate outputs\n    - /tmp/td-email.html: Temporary email HTML body\n    - /tmp/td-digest.pdf: Generated PDF digest\n    - <workspace>/archive/tech-news-digest/: Saved digest archives\n---\n\n# Tech News Digest\n\nAutomated tech news digest system with unified data source model, quality scoring pipeline, and template-based output generation.\n\n## Quick Start\n\n1. **Configuration Setup**: Default configs are in `config/defaults/`. Copy to workspace for customization:\n   ```bash\n   mkdir -p workspace/config\n   cp config/defaults/sources.json workspace/config/tech-news-digest-sources.json\n   cp config/defaults/topics.json workspace/config/tech-news-digest-topics.json\n   ```\n\n2. **Environment Variables**: \n   - `TWITTERAPI_IO_KEY` - twitterapi.io API key (optional, preferred)\n   - `X_BEARER_TOKEN` - Twitter/X official API bearer token (optional, fallback)\n   - `TAVILY_API_KEY` - Tavily Search API key, alternative to Brave (optional)\n   - `WEB_SEARCH_BACKEND` - Web search backend: auto|brave|tavily (optional, default: auto)\n   - `BRAVE_API_KEYS` - Brave Search API keys, comma-separated for rotation (optional)\n   - `BRAVE_API_KEY` - Single Brave key fallback (optional)\n   - `GITHUB_TOKEN` - GitHub personal access token (optional, improves rate limits)\n\n3. **Generate Digest**:\n   ```bash\n   # Unified pipeline (recommended) — runs all 5 sources in parallel + merge\n   python3 scripts/run-pipeline.py \\\n     --defaults config/defaults \\\n     --config workspace/config \\\n     --hours 48 --freshness pd \\\n     --archive-dir workspace/archive/tech-news-digest/ \\\n     --output /tmp/td-merged.json --verbose --force\n   ```\n\n4. **Use Templates**: Apply Discord, email, or PDF templates to merged output\n\n## Configuration Files\n\n### `sources.json` - Unified Data Sources\n```json\n{\n  \"sources\": [\n    {\n      \"id\": \"openai-rss\",\n      \"type\": \"rss\",\n      \"name\": \"OpenAI Blog\",\n      \"url\": \"https://openai.com/blog/rss.xml\",\n      \"enabled\": true,\n      \"priority\": true,\n      \"topics\": [\"llm\", \"ai-agent\"],\n      \"note\": \"Official OpenAI updates\"\n    },\n    {\n      \"id\": \"sama-twitter\",\n      \"type\": \"twitter\", \n      \"name\": \"Sam Altman\",\n      \"handle\": \"sama\",\n      \"enabled\": true,\n      \"priority\": true,\n      \"topics\": [\"llm\", \"frontier-tech\"],\n      \"note\": \"OpenAI CEO\"\n    }\n  ]\n}\n```\n\n### `topics.json` - Enhanced Topic Definitions\n```json\n{\n  \"topics\": [\n    {\n      \"id\": \"llm\",\n      \"emoji\": \"🧠\",\n      \"label\": \"LLM / Large Models\",\n      \"description\": \"Large Language Models, foundation models, breakthroughs\",\n      \"search\": {\n        \"queries\": [\"LLM latest news\", \"large language model breakthroughs\"],\n        \"must_include\": [\"LLM\", \"large language model\", \"foundation model\"],\n        \"exclude\": [\"tutorial\", \"beginner guide\"]\n      },\n      \"display\": {\n        \"max_items\": 8,\n        \"style\": \"detailed\"\n      }\n    }\n  ]\n}\n```\n\n## Scripts Pipeline\n\n### `run-pipeline.py` - Unified Pipeline (Recommended)\n```bash\npython3 scripts/run-pipeline.py \\\n  --defaults config/defaults [--config CONFIG_DIR] \\\n  --hours 48 --freshness pd \\\n  --archive-dir workspace/archive/tech-news-digest/ \\\n  --output /tmp/td-merged.json --verbose --force\n```\n- **Features**: Runs all 5 fetch steps in parallel, then merges + deduplicates + scores\n- **Output**: Final merged JSON ready for report generation (~30s total)\n- **Metadata**: Saves per-step timing and counts to `*.meta.json`\n- **GitHub Auth**: Auto-generates GitHub App token if `$GITHUB_TOKEN` not set\n- **Fallback**: If this fails, run individual scripts below\n\n### Individual Scripts (Fallback)\n\n#### `fetch-rss.py` - RSS Feed Fetcher\n```bash\npython3 scripts/fetch-rss.py [--defaults DIR] [--config DIR] [--hours 48] [--output FILE] [--verbose]\n```\n- Parallel fetching (10 workers), retry with backoff, feedparser + regex fallback\n- Timeout: 30s per feed, ETag/Last-Modified caching\n\n#### `fetch-twitter.py` - Twitter/X KOL Monitor\n```bash\npython3 scripts/fetch-twitter.py [--defaults DIR] [--config DIR] [--hours 48] [--output FILE] [--backend auto|official|twitterapiio]\n```\n- Backend auto-detection: uses twitterapi.io if `TWITTERAPI_IO_KEY` set, else official X API v2 if `X_BEARER_TOKEN` set\n- Rate limit handling, engagement metrics, retry with backoff\n\n#### `fetch-web.py` - Web Search Engine\n```bash\npython3 scripts/fetch-web.py [--defaults DIR] [--config DIR] [--freshness pd] [--output FILE]\n```\n- Auto-detects Brave API rate limit: paid plans → parallel queries, free → sequential\n- Without API: generates search interface for agents\n\n#### `fetch-github.py` - GitHub Releases Monitor\n```bash\npython3 scripts/fetch-github.py [--defaults DIR] [--config DIR] [--hours 168] [--output FILE]\n```\n- Parallel fetching (10 workers), 30s timeout\n- Auth priority: `$GITHUB_TOKEN` → GitHub App auto-generate → `gh` CLI → unauthenticated (60 req/hr)\n\n#### `fetch-reddit.py` - Reddit Posts Fetcher\n```bash\npython3 scripts/fetch-reddit.py [--defaults DIR] [--config DIR] [--hours 48] [--output FILE]\n```\n- Parallel fetching (4 workers), public JSON API (no auth required)\n- 13 subreddits with score filtering\n\n#### `merge-sources.py` - Quality Scoring & Deduplication\n```bash\npython3 scripts/merge-sources.py --rss FILE --twitter FILE --web FILE --github FILE --reddit FILE\n```\n- Quality scoring, title similarity dedup (85%), previous digest penalty\n- Output: topic-grouped articles sorted by score\n\n#### `validate-config.py` - Configuration Validator\n```bash\npython3 scripts/validate-config.py [--defaults DIR] [--config DIR] [--verbose]\n```\n- JSON schema validation, topic reference checks, duplicate ID detection\n\n#### `generate-pdf.py` - PDF Report Generator\n```bash\npython3 scripts/generate-pdf.py --input report.md --output digest.pdf [--verbose]\n```\n- Converts markdown digest to styled A4 PDF with Chinese typography (Noto Sans CJK SC)\n- Emoji icons, page headers/footers, blue accent theme. Requires `weasyprint`.\n\n#### `sanitize-html.py` - Safe HTML Email Converter\n```bash\npython3 scripts/sanitize-html.py --input report.md --output email.html [--verbose]\n```\n- Converts markdown to XSS-safe HTML email with inline CSS\n- URL whitelist (http/https only), HTML-escaped text content\n\n#### `source-health.py` - Source Health Monitor\n```bash\npython3 scripts/source-health.py --rss FILE --twitter FILE --github FILE --reddit FILE --web FILE [--verbose]\n```\n- Tracks per-source success/failure history over 7 days\n- Reports unhealthy sources (>50% failure rate)\n\n#### `summarize-merged.py` - Merged Data Summary\n```bash\npython3 scripts/summarize-merged.py --input merged.json [--top N] [--topic TOPIC]\n```\n- Human-readable summary of merged data for LLM consumption\n- Shows top articles per topic with scores and metrics\n\n## User Customization\n\n### Workspace Configuration Override\nPlace custom configs in `workspace/config/` to override defaults:\n\n- **Sources**: Append new sources, disable defaults with `\"enabled\": false`\n- **Topics**: Override topic definitions, search queries, display settings\n- **Merge Logic**: \n  - Sources with same `id` → user version takes precedence\n  - Sources with new `id` → appended to defaults\n  - Topics with same `id` → user version completely replaces default\n\n### Example Workspace Override\n```json\n// workspace/config/tech-news-digest-sources.json\n{\n  \"sources\": [\n    {\n      \"id\": \"simonwillison-rss\",\n      \"enabled\": false,\n      \"note\": \"Disabled: too noisy for my use case\"\n    },\n    {\n      \"id\": \"my-custom-blog\", \n      \"type\": \"rss\",\n      \"name\": \"My Custom Tech Blog\",\n      \"url\": \"https://myblog.com/rss\",\n      \"enabled\": true,\n      \"priority\": true,\n      \"topics\": [\"frontier-tech\"]\n    }\n  ]\n}\n```\n\n## Templates & Output\n\n### Discord Template (`references/templates/discord.md`)\n- Bullet list format with link suppression (`<link>`)\n- Mobile-optimized, emoji headers\n- 2000 character limit awareness\n\n### Email Template (`references/templates/email.md`) \n- Rich metadata, technical stats, archive links\n- Executive summary, top articles section\n- HTML-compatible formatting\n\n### PDF Template (`references/templates/pdf.md`)\n- A4 layout with Noto Sans CJK SC font for Chinese support\n- Emoji icons, page headers/footers with page numbers\n- Generated via `scripts/generate-pdf.py` (requires `weasyprint`)\n\n## Default Sources (138 total)\n\n- **RSS Feeds (49)**: AI labs, tech blogs, crypto news, Chinese tech media\n- **Twitter/X KOLs (48)**: AI researchers, crypto leaders, tech executives\n- **GitHub Repos (28)**: Major open-source projects (LangChain, vLLM, DeepSeek, Llama, etc.)\n- **Reddit (13)**: r/MachineLearning, r/LocalLLaMA, r/CryptoCurrency, r/ChatGPT, r/OpenAI, etc.\n- **Web Search (4 topics)**: LLM, AI Agent, Crypto, Frontier Tech\n\nAll sources pre-configured with appropriate topic tags and priority levels.\n\n## Dependencies\n\n```bash\npip install -r requirements.txt\n```\n\n**Optional but Recommended**:\n- `feedparser>=6.0.0` - Better RSS parsing (fallback to regex if unavailable)\n- `jsonschema>=4.0.0` - Configuration validation\n\n**All scripts work with Python 3.8+ standard library only.**\n\n## Monitoring & Operations\n\n### Health Checks\n```bash\n# Validate configuration\npython3 scripts/validate-config.py --verbose\n\n# Test RSS feeds\npython3 scripts/fetch-rss.py --hours 1 --verbose\n\n# Check Twitter API\npython3 scripts/fetch-twitter.py --hours 1 --verbose\n```\n\n### Archive Management\n- Digests automatically archived to `<workspace>/archive/tech-news-digest/`\n- Previous digest titles used for duplicate detection\n- Old archives cleaned automatically (90+ days)\n\n### Error Handling\n- **Network Failures**: Retry with exponential backoff\n- **Rate Limits**: Automatic retry with appropriate delays\n- **Invalid Content**: Graceful degradation, detailed logging\n- **Configuration Errors**: Schema validation with helpful messages\n\n## API Keys & Environment\n\nSet in `~/.zshenv` or similar:\n```bash\n# Twitter (at least one required for Twitter source)\nexport TWITTERAPI_IO_KEY=\"your_key\"        # twitterapi.io key (preferred)\nexport X_BEARER_TOKEN=\"your_bearer_token\"  # Official X API v2 (fallback)\nexport TWITTER_API_BACKEND=\"auto\"          # auto|twitterapiio|official (default: auto)\n\n# Web Search (optional, enables web search layer)\nexport WEB_SEARCH_BACKEND=\"auto\"          # auto|brave|tavily (default: auto)\nexport TAVILY_API_KEY=\"tvly-xxx\"           # Tavily Search API (free 1000/mo)\n\n# Brave Search (alternative)\nexport BRAVE_API_KEYS=\"key1,key2,key3\"     # Multiple keys, comma-separated rotation\nexport BRAVE_API_KEY=\"key1\"                # Single key fallback\nexport BRAVE_PLAN=\"free\"                   # Override rate limit detection: free|pro\n\n# GitHub (optional, improves rate limits)\nexport GITHUB_TOKEN=\"ghp_xxx\"              # PAT (simplest)\nexport GH_APP_ID=\"12345\"                   # Or use GitHub App for auto-token\nexport GH_APP_INSTALL_ID=\"67890\"\nexport GH_APP_KEY_FILE=\"/path/to/key.pem\"\n```\n\n- **Twitter**: `TWITTERAPI_IO_KEY` preferred ($3-5/mo); `X_BEARER_TOKEN` as fallback; `auto` mode tries twitterapiio first\n- **Web Search**: Tavily (preferred in auto mode) or Brave; optional, fallback to agent web_search if unavailable\n- **GitHub**: Auto-generates token from GitHub App if PAT not set; unauthenticated fallback (60 req/hr)\n- **Reddit**: No API key needed (uses public JSON API)\n\n## Cron / Scheduled Task Integration\n\n### OpenClaw Cron (Recommended)\n\nThe cron prompt should **NOT** hardcode the pipeline steps. Instead, reference `references/digest-prompt.md` and only pass configuration parameters. This ensures the pipeline logic stays in the skill repo and is consistent across all installations.\n\n#### Daily Digest Cron Prompt\n```\nRead <SKILL_DIR>/references/digest-prompt.md and follow the complete workflow to generate a daily digest.\n\nReplace placeholders with:\n- MODE = daily\n- TIME_WINDOW = past 1-2 days\n- FRESHNESS = pd\n- RSS_HOURS = 48\n- ITEMS_PER_SECTION = 3-5\n- BLOG_PICKS_COUNT = 2-3\n- EXTRA_SECTIONS = (none)\n- SUBJECT = Daily Tech Digest - YYYY-MM-DD\n- WORKSPACE = <your workspace path>\n- SKILL_DIR = <your skill install path>\n- DISCORD_CHANNEL_ID = <your channel id>\n- EMAIL = (optional)\n- LANGUAGE = English\n- TEMPLATE = discord\n\nFollow every step in the prompt template strictly. Do not skip any steps.\n```\n\n#### Weekly Digest Cron Prompt\n```\nRead <SKILL_DIR>/references/digest-prompt.md and follow the complete workflow to generate a weekly digest.\n\nReplace placeholders with:\n- MODE = weekly\n- TIME_WINDOW = past 7 days\n- FRESHNESS = pw\n- RSS_HOURS = 168\n- ITEMS_PER_SECTION = 5-8\n- BLOG_PICKS_COUNT = 3-5\n- EXTRA_SECTIONS = 📊 Weekly Trend Summary (2-3 sentences summarizing macro trends)\n- SUBJECT = Weekly Tech Digest - YYYY-MM-DD\n- WORKSPACE = <your workspace path>\n- SKILL_DIR = <your skill install path>\n- DISCORD_CHANNEL_ID = <your channel id>\n- EMAIL = (optional)\n- LANGUAGE = English\n- TEMPLATE = discord\n\nFollow every step in the prompt template strictly. Do not skip any steps.\n```\n\n#### Why This Pattern?\n- **Single source of truth**: Pipeline logic lives in `digest-prompt.md`, not scattered across cron configs\n- **Portable**: Same skill on different OpenClaw instances, just change paths and channel IDs\n- **Maintainable**: Update the skill → all cron jobs pick up changes automatically\n- **Anti-pattern**: Do NOT copy pipeline steps into the cron prompt — it will drift out of sync\n\n#### Multi-Channel Delivery Limitation\nOpenClaw enforces **cross-provider isolation**: a single session can only send messages to one provider (e.g., Discord OR Telegram, not both). If you need to deliver digests to multiple platforms, create **separate cron jobs** for each provider:\n\n```\n# Job 1: Discord + Email\n- DISCORD_CHANNEL_ID = <your-discord-channel-id>\n- EMAIL = user@example.com\n- TEMPLATE = discord\n\n# Job 2: Telegram DM\n- DISCORD_CHANNEL_ID = (none)\n- EMAIL = (none)\n- TEMPLATE = telegram\n```\nReplace `DISCORD_CHANNEL_ID` delivery with the target platform's delivery in the second job's prompt.\n\nThis is a security feature, not a bug — it prevents accidental cross-context data leakage.\n\n## Security Notes\n\n### Execution Model\nThis skill uses a **prompt template pattern**: the agent reads `digest-prompt.md` and follows its instructions. This is the standard OpenClaw skill execution model — the agent interprets structured instructions from skill-provided files. All instructions are shipped with the skill bundle and can be audited before installation.\n\n### Network Access\nThe Python scripts make outbound requests to:\n- RSS feed URLs (configured in `tech-news-digest-sources.json`)\n- Twitter/X API (`api.x.com` or `api.twitterapi.io`)\n- Brave Search API (`api.search.brave.com`)\n- Tavily Search API (`api.tavily.com`)\n- GitHub API (`api.github.com`)\n- Reddit JSON API (`reddit.com`)\n\nNo data is sent to any other endpoints. All API keys are read from environment variables declared in the skill metadata.\n\n### Shell Safety\nEmail delivery uses `send-email.py` which constructs proper MIME multipart messages with HTML body + optional PDF attachment. Subject formats are hardcoded (`Daily Tech Digest - YYYY-MM-DD`). PDF generation uses `generate-pdf.py` via `weasyprint`. The prompt template explicitly prohibits interpolating untrusted content (article titles, tweet text, etc.) into shell arguments. Email addresses and subjects must be static placeholder values only.\n\n### File Access\nScripts read from `config/` and write to `workspace/archive/`. No files outside the workspace are accessed.\n\n## Support & Troubleshooting\n\n### Common Issues\n1. **RSS feeds failing**: Check network connectivity, use `--verbose` for details\n2. **Twitter rate limits**: Reduce sources or increase interval\n3. **Configuration errors**: Run `validate-config.py` for specific issues\n4. **No articles found**: Check time window (`--hours`) and source enablement\n\n### Debug Mode\nAll scripts support `--verbose` flag for detailed logging and troubleshooting.\n\n### Performance Tuning\n- **Parallel Workers**: Adjust `MAX_WORKERS` in scripts for your system\n- **Timeout Settings**: Increase `TIMEOUT` for slow networks\n- **Article Limits**: Adjust `MAX_ARTICLES_PER_FEED` based on needs\n## Security Considerations\n\n### Shell Execution\nThe digest prompt instructs agents to run Python scripts via shell commands. All script paths and arguments are skill-defined constants — no user input is interpolated into commands. Two scripts use `subprocess`:\n- `run-pipeline.py` orchestrates child fetch scripts (all within `scripts/` directory)\n- `fetch-github.py` has two subprocess calls:\n  1. `openssl dgst -sha256 -sign` for JWT signing (only if `GH_APP_*` env vars are set — signs a self-constructed JWT payload, no user content involved)\n  2. `gh auth token` CLI fallback (only if `gh` is installed — reads from gh's own credential store)\n\nNo user-supplied or fetched content is ever interpolated into subprocess arguments. Email delivery uses `send-email.py` which builds MIME messages programmatically — no shell interpolation. PDF generation uses `generate-pdf.py` via `weasyprint`. Email subjects are static format strings only — never constructed from fetched data.\n\n### Credential & File Access\nScripts do **not** directly read `~/.config/`, `~/.ssh/`, or any credential files. All API tokens are read from environment variables declared in the skill metadata. The GitHub auth cascade is:\n1. `$GITHUB_TOKEN` env var (you control what to provide)\n2. GitHub App token generation (only if you set `GH_APP_ID`, `GH_APP_INSTALL_ID`, and `GH_APP_KEY_FILE` — uses inline JWT signing via `openssl` CLI, no external scripts involved)\n3. `gh auth token` CLI (delegates to gh's own secure credential store)\n4. Unauthenticated (60 req/hr, safe fallback)\n\nIf you prefer no automatic credential discovery, simply set `$GITHUB_TOKEN` and the script will use it directly without attempting steps 2-3.\n\n### Dependency Installation\nThis skill does **not** install any packages. `requirements.txt` lists optional dependencies (`feedparser`, `jsonschema`) for reference only. All scripts work with Python 3.8+ standard library. Users should install optional deps in a virtualenv if desired — the skill never runs `pip install`.\n\n### Input Sanitization\n- URL resolution rejects non-HTTP(S) schemes (javascript:, data:, etc.)\n- RSS fallback parsing uses simple, non-backtracking regex patterns (no ReDoS risk)\n- All fetched content is treated as untrusted data for display only\n\n### Network Access\nScripts make outbound HTTP requests to configured RSS feeds, Twitter API, GitHub API, Reddit JSON API, Brave Search API, and Tavily Search API. No inbound connections or listeners are created.\n\nFile v3.11.0:README.md\n\n# Tech News Digest\n\n> Automated tech news digest — 138 sources, 5-layer pipeline, one chat message to install.\n\n**English** | [中文](README_CN.md)\n\n[![Tests](https://github.com/draco-agent/tech-news-digest/actions/workflows/test.yml/badge.svg)](https://github.com/draco-agent/tech-news-digest/actions/workflows/test.yml)\n[![Python 3.8+](https://img.shields.io/badge/python-3.8+-blue.svg)](https://www.python.org/downloads/)\n[![ClawHub](https://img.shields.io/badge/ClawHub-tech--news--digest-blueviolet)](https://clawhub.com/draco-agent/tech-news-digest)\n[![MIT License](https://img.shields.io/badge/license-MIT-green.svg)](LICENSE)\n\n## 💬 Install in One Message\n\nTell your [OpenClaw](https://openclaw.ai) AI assistant:\n\n> **\"Install tech-news-digest and send a daily digest to #tech-news every morning at 9am\"**\n\nThat's it. Your bot handles installation, configuration, scheduling, and delivery — all through conversation.\n\nMore examples:\n\n> 🗣️ \"Set up a weekly AI digest, only LLM and AI Agent topics, deliver to Discord #ai-weekly every Monday\"\n\n> 🗣️ \"Install tech-news-digest, add my RSS feeds, and send crypto news to Telegram\"\n\n> 🗣️ \"Give me a tech digest right now, skip Twitter sources\"\n\nOr install via CLI:\n```bash\nclawhub install tech-news-digest\n```\n\n## 📊 What You Get\n\nA quality-scored, deduplicated tech digest built from **138 sources**:\n\n| Layer | Sources | What |\n|-------|---------|------|\n| 📡 RSS | 49 feeds | OpenAI, Anthropic, Ben's Bites, HN, 36氪, CoinDesk… |\n| 🐦 Twitter/X | 48 KOLs | @karpathy, @VitalikButerin, @sama, @elonmusk… |\n| 🔍 Web Search | 4 topics | Brave Search API with freshness filters |\n| 🐙 GitHub | 28 repos | Releases from key projects (LangChain, vLLM, DeepSeek, Llama…) |\n| 🗣️ Reddit | 13 subs | r/MachineLearning, r/LocalLLaMA, r/CryptoCurrency… |\n\n### Pipeline\n\n```\n       run-pipeline.py (~30s)\n              ↓\n  RSS ─┐\n  Twitter ─┤\n  Web ─────┤── parallel fetch ──→ merge-sources.py\n  GitHub ──┤\n  Reddit ──┘\n              ↓\n  Quality Scoring → Deduplication → Topic Grouping\n              ↓\n    Discord / Email / PDF output\n```\n\n**Quality scoring**: priority source (+3), multi-source cross-ref (+5), recency (+2), engagement (+1), Reddit score bonus (+1/+3/+5), already reported (-5).\n\n## ⚙️ Configuration\n\n- `config/defaults/sources.json` — 138 built-in sources\n- `config/defaults/topics.json` — 4 topics with search queries & Twitter queries\n- User overrides in `workspace/config/` take priority\n\n## 🎨 Customize Your Sources\n\nWorks out of the box with 138 built-in sources — but fully customizable. Copy the defaults to your workspace config and override:\n\n```bash\n# Copy and customize\ncp config/defaults/sources.json workspace/config/tech-news-digest-sources.json\ncp config/defaults/topics.json workspace/config/tech-news-digest-topics.json\n```\n\nYour overlay file **merges** with defaults:\n- **Override** a source by matching its `id` — your version replaces the default\n- **Add** new sources with a unique `id` — appended to the list\n- **Disable** a built-in source — set `\"enabled\": false` on the matching `id`\n\n```json\n{\n  \"sources\": [\n    {\"id\": \"my-blog\", \"type\": \"rss\", \"enabled\": true, \"url\": \"https://myblog.com/feed\", \"topics\": [\"llm\"]},\n    {\"id\": \"openai-blog\", \"enabled\": false}\n  ]\n}\n```\n\nNo need to copy the entire file — just include what you want to change.\n\n## 🔧 Optional Setup\n\nAll environment variables are optional. The pipeline runs with whatever sources are available.\n\n```bash\nexport TWITTERAPI_IO_KEY=\"...\"  # twitterapi.io (~$5/mo) — enables Twitter layer\nexport X_BEARER_TOKEN=\"...\"     # Twitter/X official API — alternative Twitter backend\nexport TAVILY_API_KEY=\"tvly-xxx\"  # Tavily Search API (alternative, free 1000/mo)\nexport BRAVE_API_KEYS=\"k1,k2,k3\" # Brave Search API keys (comma-separated, rotation)\nexport BRAVE_API_KEY=\"...\"       # Fallback: single Brave key\nexport GITHUB_TOKEN=\"...\"       # GitHub API — higher rate limits (auto-generated from GitHub App if unset)\nexport TWITTER_API_BACKEND=\"auto\" # auto|twitterapiio|official (default: auto)\nexport BRAVE_PLAN=\"free\"         # Override Brave rate limit detection: free|pro\nexport WEB_SEARCH_BACKEND=\"auto\" # auto|brave|tavily (default: auto)\npip install weasyprint           # Enables PDF report generation\n```\n\n## 📂 Repository\n\n**GitHub**: [github.com/draco-agent/tech-news-digest](https://github.com/draco-agent/tech-news-digest)\n\n## 📄 License\n\nMIT License — see [LICENSE](LICENSE) for details.\n\nFile v3.11.0:_meta.json\n\n{\n  \"ownerId\": \"kn74589cx1nbhnc3x0f3nwre39814699\",\n  \"slug\": \"tech-news-digest\",\n  \"version\": \"3.11.0\",\n  \"publishedAt\": 1772295728196\n}\n\nFile v3.11.0:references/digest-prompt.md\n\n# Digest Prompt Template\n\nReplace `<...>` placeholders before use. Daily defaults shown; weekly overrides in parentheses.\n\n## Placeholders\n\n| Placeholder | Default | Weekly Override |\n|-------------|---------|----------------|\n| `<MODE>` | `daily` | `weekly` |\n| `<TIME_WINDOW>` | `past 1-2 days` | `past 7 days` |\n| `<FRESHNESS>` | `pd` | `pw` |\n| `<RSS_HOURS>` | `48` | `168` |\n| `<ITEMS_PER_SECTION>` | `3-5` | `5-8` |\n| `<BLOG_PICKS_COUNT>` | `2-3` | `3-5` |\n| `<EXTRA_SECTIONS>` | *(none)* | `📊 Weekly Trend Summary` |\n| `<SUBJECT>` | `Daily Tech Digest - YYYY-MM-DD` | `Weekly Tech Digest - YYYY-MM-DD` |\n| `<WORKSPACE>` | Your workspace path | |\n| `<SKILL_DIR>` | Installed skill directory | |\n| `<DISCORD_CHANNEL_ID>` | Target channel ID | |\n| `<EMAIL>` | *(optional)* Recipient email | |\n| `<EMAIL_FROM>` | *(optional)* e.g. `MyBot <bot@example.com>` | |\n| `<LANGUAGE>` | `Chinese` | |\n| `<TEMPLATE>` | `discord` / `email` / `markdown` | |\n| `<DATE>` | Today's date YYYY-MM-DD (caller provides) | |\n| `<VERSION>` | Read from SKILL.md frontmatter | |\n\n---\n\nGenerate the <MODE> tech digest for **<DATE>**. Use `<DATE>` as the report date — do NOT infer it.\n\n## Configuration\n\nRead config files (workspace overrides take priority over defaults):\n1. **Sources**: `<WORKSPACE>/config/tech-news-digest-sources.json` → fallback `<SKILL_DIR>/config/defaults/sources.json`\n2. **Topics**: `<WORKSPACE>/config/tech-news-digest-topics.json` → fallback `<SKILL_DIR>/config/defaults/topics.json`\n\n## Context: Previous Report\n\nRead the most recent file from `<WORKSPACE>/archive/tech-news-digest/` to avoid repeats and follow up on developing stories. Skip if none exists.\n\n## Data Collection Pipeline\n\n**Use the unified pipeline** (runs all 5 sources in parallel, ~30s):\n\n```bash\npython3 <SKILL_DIR>/scripts/run-pipeline.py \\\n  --defaults <SKILL_DIR>/config/defaults \\\n  --config <WORKSPACE>/config \\\n  --hours <RSS_HOURS> --freshness <FRESHNESS> \\\n  --archive-dir <WORKSPACE>/archive/tech-news-digest/ \\\n  --output /tmp/td-merged.json --verbose --force\n```\n\nIf it fails, run individual scripts in `<SKILL_DIR>/scripts/` (see each script's `--help`), then merge with `merge-sources.py`.\n\n## Report Generation\n\nGet a structured overview:\n```bash\npython3 <SKILL_DIR>/scripts/summarize-merged.py --input /tmp/td-merged.json --top <ITEMS_PER_SECTION>\n```\n\nUse this output to select articles — **do NOT write ad-hoc Python to parse the JSON**. Apply the template from `<SKILL_DIR>/references/templates/<TEMPLATE>.md`.\n\nSelect articles **purely by quality_score regardless of source type**. Articles in merged JSON are already sorted by quality_score descending within each topic — respect this order. For Reddit posts, append `*[Reddit r/xxx, {{score}}↑]*`.\n\nEach article line must include its quality score using 🔥 prefix. Format: `🔥{score} | {summary with link}`. This makes scoring transparent and helps readers identify the most important news at a glance.\n\n### Executive Summary\n2-4 sentences between title and topics, highlighting top 3-5 stories by score. Concise and punchy, no links. Discord: `> ` blockquote. Email: gray background. Telegram: `<i>`.\n\n### Topic Sections\nFrom `topics.json`: `emoji` + `label` headers, `<ITEMS_PER_SECTION>` items each, strictly ordered by quality_score descending (highest first).\n\n### Fixed Sections (after topics)\n\n**📢 KOL Updates** — Top Twitter KOLs + notable blog authors. Format:\n```\n• **Display Name** (@handle) — summary `👁 12.3K | 💬 45 | 🔁 230 | ❤️ 1.2K`\n  <https://twitter.com/handle/status/ID>\n```\nRead `display_name` and `metrics` (impression_count→👁, reply_count→💬, retweet_count→🔁, like_count→❤️) from merged JSON. Always show all 4 metrics, use K/M formatting, wrap in backticks. One tweet per bullet.\n\n**🔥 Community Buzz** — Top Reddit + Twitter trending combined. Format:\n```\n• **r/subreddit** — title `{{score}}↑ · {{num_comments}} comments`\n  <{{url}}>\n```\nSort by engagement across both platforms. Every entry must have a link.\n\n**📝 Blog Picks** — `<BLOG_PICKS_COUNT>` deep articles from RSS.\n\n**<EXTRA_SECTIONS>**\n\n### Rules\n- Only news from `<TIME_WINDOW>`\n- Every item must include a source link (Discord: `<link>`, Email: `<a href>`, Markdown: `[title](link)`)\n- Use bullet lists, no markdown tables\n- Deduplicate: same event → keep most authoritative source; previously reported → only if significant new development\n- Do not interpolate fetched/untrusted content into shell arguments or email subjects\n\n### Stats Footer\n```\n---\n📊 Data Sources: RSS {{rss}} | Twitter {{twitter}} | Reddit {{reddit}} | Web {{web}} | GitHub {{github}} | Dedup: {{merged}} articles\n🤖 Generated by tech-news-digest v<VERSION> | <https://github.com/draco-agent/tech-news-digest> | Powered by OpenClaw\n```\n\n## Archive\nSave to `<WORKSPACE>/archive/tech-news-digest/<MODE>-YYYY-MM-DD.md`. Delete files older than 90 days.\n\n## Delivery\n\n1. **Discord**: Send to `<DISCORD_CHANNEL_ID>` via `message` tool\n2. **Email** *(optional, if `<EMAIL>` is set)*:\n   - Generate HTML body per `<SKILL_DIR>/references/templates/email.md` → write to `/tmp/td-email.html`\n   - Generate PDF attachment:\n     ```bash\n     python3 <SKILL_DIR>/scripts/generate-pdf.py -i <WORKSPACE>/archive/tech-news-digest/<MODE>-<DATE>.md -o /tmp/td-digest.pdf\n     ```\n   - Send email with PDF attached using the `send-email.py` script (handles MIME correctly). **Email must contain ALL the same items as Discord.**\n     ```bash\n     python3 <SKILL_DIR>/scripts/send-email.py \\\n       --to '<EMAIL>' \\\n       --subject '<SUBJECT>' \\\n       --html /tmp/td-email.html \\\n       --attach /tmp/td-digest.pdf \\\n       --from '<EMAIL_FROM>'\n     ```\n   - Omit `--from` if `<EMAIL_FROM>` is not set. Omit `--attach` if PDF generation failed. SUBJECT must be a static string. If delivery fails, log error and continue.\n\nWrite the report in <LANGUAGE>.\n\nFile v3.11.0:references/templates/discord.md\n\n# Tech Digest Discord Template\n\nDiscord-optimized format with bullet points and link suppression.\n\n## Template Structure\n\n```markdown\n# 🚀 Tech Digest - {{DATE}}\n\n{{#topics}}\n## {{emoji}} {{label}}\n\n{{#articles}}\n• 🔥{{quality_score}} | {{title}}\n  <{{link}}>\n  {{#multi_source}}*[{{source_count}} sources]*{{/multi_source}}\n\n{{/articles}}\n{{/topics}}\n\n---\n📊 Data Sources: RSS {{rss_count}} | Twitter {{twitter_count}} | Reddit {{reddit_count}} | Web {{web_count}} | GitHub {{github_count}} releases | After dedup: {{merged_count}} articles\n🤖 Generated by tech-news-digest v{{version}} | <https://github.com/draco-agent/tech-news-digest> | Powered by OpenClaw\n```\n\n## Delivery\n\n- **Default: Channel** — Send to the Discord channel specified by `DISCORD_CHANNEL_ID`\n- Use `message` tool with `target` set to the channel ID for channel delivery\n- For DM delivery instead, set `target` to a user ID\n\n## Discord-Specific Features\n\n- **Link suppression**: Wrap links in `<>` to prevent embeds\n- **Bullet format**: Use `•` for clean mobile display  \n- **No tables**: Discord mobile doesn't handle markdown tables well\n- **Emoji headers**: Visual hierarchy with topic emojis\n- **Concise metadata**: Source count and multi-source indicators\n- **Character limits**: Discord messages have 2000 char limit, may need splitting\n\n## Example Output\n\n```markdown\n# 🚀 Tech Digest - 2026-02-15\n\n## 🧠 LLM / Large Models\n\n• 🔥15 | OpenAI releases GPT-5 with breakthrough reasoning capabilities\n  <https://openai.com/blog/gpt5-announcement>\n  *[3 sources]*\n\n• 🔥12 | Meta's Llama 3.1 achieves new MMLU benchmarks\n  <https://ai.meta.com/blog/llama-31-release>\n\n## 🤖 AI Agent\n\n• 🔥14 | LangChain launches production-ready agent framework\n  <https://blog.langchain.dev/production-agents>\n\n## 💰 Cryptocurrency\n\n• 🔥18 | Bitcoin reaches new ATH at $67,000 amid ETF approval\n  <https://coindesk.com/markets/btc-ath-etf>\n  *[2 sources]*\n\n## 📢 KOL Updates\n\n• **Elon Musk** (@elonmusk) — Confirmed X's crypto trading feature `👁 2.1M | 💬 12.3K | 🔁 8.5K | ❤️ 49.8K`\n  <https://twitter.com/elonmusk/status/123456789>\n• **@saylor** — Valentine's BTC enthusiasm `👁 450K | 💬 1.2K | 🔁 3.1K | ❤️ 13K`\n  <https://twitter.com/saylor/status/987654321>\n\n---\n📊 Data Sources: RSS 285 | Twitter 67 | Reddit 45 | Web 60 | GitHub 29 releases | After dedup: 95 articles\n```\n\n## Variables\n\n- `{{DATE}}` - Report date (YYYY-MM-DD format)\n- `{{topics}}` - Array of topic objects\n- `{{emoji}}` - Topic emoji \n- `{{label}}` - Topic display name\n- `{{articles}}` - Array of article objects per topic\n- `{{title}}` - Article title (truncated if needed)\n- `{{link}}` - Article URL\n- `{{quality_score}}` - Article quality score (higher = more important)\n- `{{multi_source}}` - Boolean, true if article from multiple sources\n- `{{source_count}}` - Number of sources for this article\n- `{{total_sources}}` - Total number of sources used\n- `{{total_articles}}` - Total articles in digest\n\nFile v3.11.0:references/templates/email.md\n\n# Tech Digest Email Template\n\nHTML email format optimized for Gmail/Outlook rendering.\n\n## Delivery\n\nSend via `gog gmail send` with `--body-html` flag:\n```bash\ngog gmail send --to '<EMAIL>' --subject '<SUBJECT>' --body-html '<HTML_CONTENT>'\n```\n\n**Important**: Use `--body-html`, NOT `--body`. Plain text markdown will not render properly in email clients.\n\n## Template Structure\n\nThe agent should generate an HTML email body. Use inline styles (email clients strip `<style>` blocks).\n\n```html\n<div style=\"max-width:640px;margin:0 auto;font-family:-apple-system,BlinkMacSystemFont,'Segoe UI',Roboto,sans-serif;color:#1a1a1a;line-height:1.6\">\n\n  <h1 style=\"font-size:22px;border-bottom:2px solid #e5e5e5;padding-bottom:8px\">\n    🐉 {{TITLE}}\n  </h1>\n\n  <!-- Optional: Executive Summary for weekly -->\n  <p style=\"color:#555;font-size:14px;background:#f8f9fa;padding:12px;border-radius:6px\">\n    {{SUMMARY}}\n  </p>\n\n  <!-- Topic Section -->\n  <h2 style=\"font-size:17px;margin-top:24px;color:#333\">{{emoji}} {{label}}</h2>\n  <ul style=\"padding-left:20px\">\n    <li style=\"margin-bottom:10px\">\n      <strong>🔥{{quality_score}}</strong> {{title}} — {{description}}\n      <br><a href=\"{{link}}\" style=\"color:#0969da;font-size:13px\">{{link}}</a>\n    </li>\n  </ul>\n\n  <!-- Repeat for each topic -->\n\n  <!-- KOL Section: Read metrics from twitter JSON data (metrics.impression_count, reply_count, retweet_count, like_count). One tweet per <li>. -->\n  <h2 style=\"font-size:17px;margin-top:24px;color:#333\">📢 KOL Updates</h2>\n  <ul style=\"padding-left:20px\">\n    <li style=\"margin-bottom:10px\">\n      <strong>{{display_name}}</strong> (@{{handle}}) — {{summary}}\n      <br><code style=\"font-size:12px;color:#888;background:#f4f4f4;padding:2px 6px;border-radius:3px\">👁 {{views}} | 💬 {{replies}} | 🔁 {{retweets}} | ❤️ {{likes}}</code>\n      <br><a href=\"{{tweet_link}}\" style=\"color:#0969da;font-size:13px\">{{tweet_link}}</a>\n    </li>\n  </ul>\n\n  <!-- Twitter/X Trending Section: Each entry must include at least one reference link -->\n  <h2 style=\"font-size:17px;margin-top:24px;color:#333\">🔥 Community Buzz</h2>\n  <ul style=\"padding-left:20px\">\n    <li style=\"margin-bottom:10px\">\n      <strong>{{trending_topic}}</strong> — {{summary}}\n      <br><a href=\"{{reference_link}}\" style=\"color:#0969da;font-size:13px\">{{reference_link}}</a>\n    </li>\n  </ul>\n\n  <!-- Blog / Releases sections -->\n\n  <!-- Footer -->\n  <hr style=\"border:none;border-top:1px solid #e5e5e5;margin:24px 0\">\n  <p style=\"font-size:12px;color:#888\">\n    📊 Data Sources: RSS {{rss_count}} | Twitter {{twitter_count}} | Reddit {{reddit_count}} | Web {{web_count}} | GitHub {{github_count}} releases | After dedup: {{merged_count}} articles\n    <br>🤖 Generated by <a href=\"https://github.com/draco-agent/tech-news-digest\" style=\"color:#0969da\">tech-news-digest</a> v{{version}} | Powered by <a href=\"https://openclaw.ai\" style=\"color:#0969da\">OpenClaw</a>\n  </p>\n\n</div>\n```\n\n## Style Guidelines\n\n- **Max width**: 640px centered (mobile-friendly)\n- **Fonts**: System font stack (no web fonts in email)\n- **All styles inline**: Email clients strip `<style>` tags\n- **Links**: Use full URLs, styled with `color:#0969da`\n- **Headings**: h1 for title (22px), h2 for topics (17px)\n- **Lists**: `<ul>` with `<li>`, adequate spacing\n- **Footer**: Small gray text with stats\n- **No images**: Pure text/HTML for maximum compatibility\n- **No tables for layout**: Use div + inline styles\n\n## Example Output\n\n```html\n<div style=\"max-width:640px;margin:0 auto;font-family:-apple-system,BlinkMacSystemFont,'Segoe UI',Roboto,sans-serif;color:#1a1a1a;line-height:1.6\">\n\n  <h1 style=\"font-size:22px;border-bottom:2px solid #e5e5e5;padding-bottom:8px\">\n    🐉 Daily Tech Digest — 2026-02-15\n  </h1>\n\n  <h2 style=\"font-size:17px;margin-top:24px;color:#333\">🧠 LLM / Large Models</h2>\n  <ul style=\"padding-left:20px\">\n    <li style=\"margin-bottom:10px\">\n      <strong>GPT-5.2 achieves first theoretical physics discovery</strong> — Collaboration with IAS, Cambridge, Harvard on gluon interactions\n      <br><a href=\"https://twitter.com/OpenAI/status/2022390096625078389\" style=\"color:#0969da;font-size:13px\">twitter.com/OpenAI</a>\n    </li>\n    <li style=\"margin-bottom:10px\">\n      <strong>ByteDance releases Doubao 2.0</strong> — Full upgrade across Agent, image, and video\n      <br><a href=\"https://www.jiqizhixin.com/articles/2026-02-14-9\" style=\"color:#0969da;font-size:13px\">jiqizhixin.com</a>\n    </li>\n    <li style=\"margin-bottom:10px\">\n      <strong>Dario Amodei: nearing the end of exponential growth</strong> — In-depth Anthropic CEO interview\n      <br><a href=\"https://www.dwarkesh.com/p/dario-amodei-2\" style=\"color:#0969da;font-size:13px\">dwarkesh.com</a>\n    </li>\n  </ul>\n\n  <h2 style=\"font-size:17px;margin-top:24px;color:#333\">🤖 AI Agent</h2>\n  <ul style=\"padding-left:20px\">\n    <li style=\"margin-bottom:10px\">\n      <strong>Stanford AI Town startup raises $100M</strong> — Backed by Fei-Fei Li, Karpathy\n      <br><a href=\"https://www.qbitai.com/2026/02/380347.html\" style=\"color:#0969da;font-size:13px\">qbitai.com</a>\n    </li>\n  </ul>\n\n  <h2 style=\"font-size:17px;margin-top:24px;color:#333\">💰 Cryptocurrency</h2>\n  <ul style=\"padding-left:20px\">\n    <li style=\"margin-bottom:10px\">\n      <strong>X to launch crypto & stock trading</strong> — Smart Cashtags feature coming soon\n      <br><a href=\"https://www.theblock.co/post/389952\" style=\"color:#0969da;font-size:13px\">theblock.co</a>\n    </li>\n  </ul>\n\n  <hr style=\"border:none;border-top:1px solid #e5e5e5;margin:24px 0\">\n  <p style=\"font-size:12px;color:#888\">\n    📊 Data Sources: RSS 287 | Twitter 71 | Reddit 45 | Web 60 | GitHub 29 releases | After dedup: 140 articles\n    <br>Generated by Tech News Digest\n  </p>\n\n</div>\n```\n\nFile v3.11.0:references/templates/pdf.md\n\n# Tech Digest PDF Template\n\nProfessional PDF output with Chinese typography, emoji icons, and A4 layout.\n\n## Generation\n\nGenerate PDF from the markdown report using `generate-pdf.py`:\n\n```bash\npython3 scripts/generate-pdf.py --input /tmp/td-report.md --output /tmp/td-digest.pdf\n```\n\n## Prerequisites\n\n- **weasyprint**: `pip install weasyprint`\n- **Chinese fonts**: `apt install fonts-noto-cjk` (Noto Sans CJK SC)\n\n## Workflow\n\n1. Generate the digest in **markdown format** first (same as Discord template output)\n2. Save the markdown to a temp file (e.g., `/tmp/td-report.md`)\n3. Run `generate-pdf.py` to convert to PDF\n4. Optionally attach the PDF to Discord or email\n\n## Features\n\n- **A4 layout** with 2cm/2.5cm margins\n- **Noto Sans CJK SC** font for native Chinese rendering\n- **Emoji support** — section icons (🧠🤖💰🔬) render correctly\n- **Page headers/footers** — \"Tech Digest\" header, page numbers\n- **Blue accent color scheme** — headers, links, blockquote borders\n- **Blockquote summary** — highlighted executive summary area\n- **Source links** — compact, below each item\n- **Responsive bullet lists** — clean indentation\n\n## Example Markdown Input\n\nThe PDF generator accepts the same markdown format as the Discord template:\n\n```markdown\n# 🚀 科技日报 - 2026-02-25\n\n> 今日要闻：OpenAI 发布新模型，Anthropic 推出 Claude 4...\n\n## 🧠 LLM / 大语言模型\n\n• **OpenAI 发布 GPT-5** — 全新推理能力突破\n  <https://openai.com/blog/gpt5>\n\n• **Anthropic Claude 4 上线** — 更强的代码能力\n  <https://anthropic.com/claude-4>\n\n## 💰 Crypto / 区块链\n\n• **以太坊 Pectra 升级完成** — EIP-7702 正式上线\n  <https://ethereum.org/pectra>\n\n---\n📊 数据源: RSS 180 | Twitter 98 | Reddit 45 | Web 20 | GitHub 15\n🤖 Generated by tech-news-digest v3.9.1\n```\n\n## Delivery\n\n```bash\n# Generate PDF\npython3 scripts/generate-pdf.py -i /tmp/td-report.md -o /tmp/td-digest.pdf\n\n# Attach to Discord\n# (use message tool with filePath parameter)\n\n# Attach to email\nmail -a /tmp/td-digest.pdf -s \"Tech Digest\" recipient@example.com < /dev/null\n```\n\nFile v3.11.0:CHANGELOG.md\n\n# Changelog\n\n## v3.11.0\n\n- **Tavily Search backend**: Alternative to Brave Search via `TAVILY_API_KEY` + `WEB_SEARCH_BACKEND` env\n- **Quality scores in output**: 🔥 score prefix on every article, strict descending order per topic\n- **Domain limit fix**: Exempt x.com/github.com/reddit.com from per-topic domain limits (#1)\n- **Brave multi-key**: `BRAVE_API_KEYS` for comma-separated key rotation\n- **Config naming**: User overlay files renamed to `tech-news-digest-sources.json` / `tech-news-digest-topics.json`\n- **Tests**: 41 unit + integration tests with real fixture data, GitHub Actions CI (Python 3.9 + 3.12)\n- **Docs**: Full env var alignment, Network Access/Shell Safety updates, README badges, CN sync\n\n## v3.10.3\n\n- **Docs**: Align API Keys & Environment with all 10 actual env vars\n- **Docs**: Update Network Access (add Reddit) and Shell Safety (send-email.py + generate-pdf.py)\n- **Refactor**: Rename user overlay configs to `tech-news-digest-sources.json` / `tech-news-digest-topics.json` to avoid naming conflicts\n\n## v3.10.2\n\n- **Fix domain limits**: Exempt multi-author platforms (x.com, github.com, reddit.com) from per-topic domain limits — previously 77 tweets were cut to 12 (#1)\n- **Brave multi-key**: Prefer `BRAVE_API_KEYS` (comma-separated) over `BRAVE_API_KEY` for key rotation in `fetch-web.py`\n\n## v3.10.1\n\n- **Fix email MIME**: New `send-email.py` — proper multipart MIME construction for HTML body + PDF attachment (replaces broken `mail -a -A` approach)\n- **Docs alignment**: README + SKILL.md updated to v3.10 (source counts, PDF, all scripts documented)\n\n## v3.10.0\n\n- **PDF generation**: New `generate-pdf.py` script — converts markdown digest to styled A4 PDF with Chinese typography (Noto Sans CJK SC), emoji icons, page headers/footers, blue accent theme. Requires `weasyprint`.\n- **PDF template**: `references/templates/pdf.md` with usage docs and examples\n\n## v3.9.1\n\n- Remove unused markdown and telegram templates\n- Add `sanitize-html.py` for safe markdown→HTML email conversion (XSS-safe, inline CSS)\n\n## v3.9.0\n\n- **URL-based dedup**: merge-sources now deduplicates by normalized URL (domain+path) before title similarity, catching cross-source duplicates\n- **Brave rate limit caching**: `detect_brave_rate_limit()` results cached for 24h; supports `BRAVE_PLAN=free|pro` env override\n- **source-health**: Now tracks Reddit (`--reddit`) and Web (`--web`) sources; flexible key detection\n- **run-pipeline**: `--skip` (comma-separated step names) and `--reuse-dir` (reuse intermediate outputs) for partial reruns\n\n## v3.8.1\n\n- **merge-sources**: Fix `getattr` → direct `args.reddit`; domain limit stats now accurate; SequenceMatcher early-exit for >30% length diff\n- **merge-sources**: RSS priority sources get +2 extra score bonus (prevent drowning by low-engagement tweets)\n- **run-pipeline**: Add `--twitter-backend` parameter (transparent passthrough); clean up tmp dir after success\n- **fetch-rss**: Warn when feedparser not installed (basic regex parser fallback)\n- **config_loader**: Validate required fields (id, type, enabled) on source load, skip invalid with warning\n\n## v3.8.0\n\n- **twitterapiio pagination**: Fetches up to 2 pages (40 tweets) for high-volume users; logs truncation warning\n- **Unified tweet limit**: `MAX_TWEETS_PER_USER` 10→20 for official backend (matches twitterapiio)\n- **Shared result helpers**: `_make_result()` / `_make_error()` on base class, reduces duplication\n- **Smarter rate limiting**: `RateLimiter` class with `threading.Lock` for twitterapiio (5 QPS); replaces per-thread sleep\n- **Retry improvements**: `RETRY_COUNT` 1→2 (3 attempts); twitterapiio 429 wait 60s→5s\n- **Tweet text limit**: 200→280 chars (matches Twitter's actual limit)\n- **Empty result format**: Now matches normal output structure for consistent downstream parsing\n- **Removed redundant isReply filter** in twitterapiio (API already excludes replies)\n\n## v3.7.1\n\n- **twitterapi.io bugfix**: Fix response envelope parsing (`data.tweets` not top-level `tweets`)\n- **twitterapi.io concurrency**: 3-worker parallel fetch with progress logs showing tweet counts and top likes\n- **test-pipeline.sh revamp**: `--only`, `--skip`, `--topics`, `--ids`, `--twitter-backend` filtering; per-step timing; detailed `--help`\n\n## v3.7.0\n\n- **twitterapi.io backend**: Alternative Twitter data source via `TWITTERAPI_IO_KEY` — no username→ID resolution needed, simpler API, same normalized output format\n- **Backend auto-detection**: `TWITTER_API_BACKEND=auto` (default) uses twitterapi.io if key is set, else falls back to official X API v2\n- **`--backend` CLI arg**: Override env var per invocation (`official`, `twitterapiio`, `auto`)\n- **Backend abstraction**: `fetch-twitter.py` refactored with `TwitterBackend` base class and two implementations (`OfficialBackend`, `TwitterApiIoBackend`)\n\n## v3.6.3\n\n- Add GitHub source: zeroclaw-labs/zeroclaw (137→138 total, 27→28 GitHub)\n\n## v3.6.2\n\n- Add 3 GitHub sources: cloudflare/moltworker, sipeed/picoclaw, HKUDS/nanobot (134→137 total, 24→27 GitHub)\n\n## v3.6.1\n\n- Prompt review & optimization pass (no functional changes)\n\n## v3.6.0\n\n- Simplify digest-prompt: 232→122 lines (-47%), remove fallback scripts block, merge redundant rules\n- Add optional `<EMAIL_FROM>` placeholder for sender display name\n- Add \"Environment vs Code\" separation rule to CONTRIBUTING.md\n\n## v3.5.1\n\n- Email delivery: prefer `mail` (msmtp) over `gog`, remove redundant fallback options\n- Require email content to match Discord (no abbreviation or skipped sections)\n- Add CONTRIBUTING.md with development conventions\n\n## v3.5.0\n\n- **Unified source count**: 134 sources (49 RSS + 48 Twitter + 24 GitHub + 13 Reddit)\n- Updated README source counts and sub-totals\n\n## v3.4.9\n\n- Declare `openssl` as optional binary in SKILL.md (used for GitHub App JWT signing)\n\n## v3.4.8\n\n- **New `summarize-merged.py` helper**: Outputs structured human-readable summary of merged data, sorted by quality score with metrics/sources\n- **Prevent ad-hoc JSON parsing**: `digest-prompt.md` now instructs agents to use `summarize-merged.py` instead of writing inline Python (which often failed with `AttributeError` on nested structures)\n\n## v3.4.7\n\n- **Inline GitHub App JWT signing**: Remove `GH_APP_TOKEN_SCRIPT` env var entirely. Token generation now built into `fetch-github.py` using `openssl` CLI for RS256 signing — no external scripts executed, no arbitrary code execution risk.\n- Only 3 env vars needed: `GH_APP_ID`, `GH_APP_INSTALL_ID`, `GH_APP_KEY_FILE`\n- Remove unused imports, fix bare excepts across all scripts\n\n## v3.4.6\n\n- Add `reddit` to config/schema.json source type enum (was missing, caused validation mismatch)\n- Rename all archive paths `tech-digest/` → `tech-news-digest/` for consistency\n- Fix Discord template: default delivery is channel (via DISCORD_CHANNEL_ID), not DM\n- GH_APP_TOKEN_SCRIPT: add trust warning in code and env var description\n- Path placeholders: SKILL.md uses `<workspace>/` consistently with digest-prompt.md\n\n## v3.4.5\n\n- Fix source count inconsistencies across docs (131/132 → 133: 49 RSS + 49 Twitter + 22 GitHub + 13 Reddit)\n- Rename legacy `tech-digest` references to `tech-news-digest` in comments, descriptions, and cache file paths\n\n## v3.4.4\n\n- Remove hardcoded Discord channel ID from SKILL.md (use `<your-discord-channel-id>` placeholder)\n- Cron prompt examples: Chinese → English, default LANGUAGE = English\n- Remove outdated \"Migration from v1.x\" section\n\n## v3.4.3\n\n- **Audit compliance**: Address all ClawHub Code Insights findings:\n  - Declare `gh` as optional binary in SKILL.md metadata\n  - Document credential access cascade and file access scope in security section\n  - Add \"Dependency Installation\" section clarifying skill never runs `pip install`\n  - Explicitly state scripts do not read `~/.config/`, `~/.ssh/`, or arbitrary credential files\n\n## v3.4.2\n\n- **Remove hardcoded GitHub App credentials**: App ID, install ID, key file path, and token script path now read exclusively from env vars (`GH_APP_ID`, `GH_APP_INSTALL_ID`, `GH_APP_KEY_FILE`, `GH_APP_TOKEN_SCRIPT`). No defaults — if not set, this auth method is silently skipped.\n- **Declare new env vars in SKILL.md**: All 4 GitHub App env vars declared in metadata\n- **Fix security docs**: Updated Shell Execution section to accurately describe `subprocess.run()` usage in `run-pipeline.py` and `fetch-github.py`\n\n## v3.4.1\n\n- **KOL Display Names**: KOL Updates section now shows \"Sam Altman (@sama)\" instead of bare \"@sama\" across all templates (Discord, Email, Telegram)\n- **`display_name` in Merged JSON**: `merge-sources.py` propagates Twitter source `name` to article-level `display_name` field, eliminating need to re-read raw Twitter data\n- **New Twitter Sources**: Added @OpenClawAI (official) and @steipete (Peter Steinberger), total 49 Twitter KOLs / 133 sources\n- **Enforce Unified Pipeline**: `digest-prompt.md` now says \"You MUST use\" `run-pipeline.py`, individual steps demoted to `<details>` fallback with `--force` flags\n\n## v3.4.0\n\n- **Unified Pipeline**: New `run-pipeline.py` runs all 5 fetch steps (RSS, Twitter, GitHub, Reddit, Web) in parallel, then merges — total ~30s vs ~3-4min sequential. Digest prompt updated to use this by default.\n- **Reddit Parallel Fetch**: `fetch-reddit.py` now uses `ThreadPoolExecutor(max_workers=4)` instead of sequential requests with `sleep(1)`\n- **Reddit 403 Fix**: Added explicit `ssl.create_default_context()` and `Accept-Language` header to fix Reddit blocking Python's default `urllib` TLS fingerprint\n- **Brave API Auto-Concurrency**: `fetch-web.py` probes `x-ratelimit-limit` header at startup — paid plans auto-switch to parallel queries, free plans stay sequential\n- **GitHub Auto-Auth**: `fetch-github.py` resolves tokens in priority order: `$GITHUB_TOKEN` → GitHub App auto-generate → `gh` CLI → unauthenticated. No manual token setup needed if GitHub App credentials exist.\n- **Timeout Increase**: All fetch scripts 15s → 30s per HTTP request; pipeline per-step subprocess 120s → 180s\n- **Pipeline Metadata**: `run-pipeline.py` saves `*.meta.json` with per-step timing, counts, and status\n\n## v3.3.2\n\n- **Declare tools and file access**: Added `tools` (python3 required, gog optional) and `files` (read/write paths) to SKILL.md metadata, addressing VirusTotal \"undeclared tools/binaries\" and \"modify workspace files\" audit findings\n- **Added `metadata.openclaw.requires`**: Declares `python3` binary dependency\n\n## v3.3.1\n\n- **Remove anthropic-rss mirror**: Removed third-party community RSS mirror (`anthropic-rss`) to eliminate supply chain risk flagged by VirusTotal Code Insights. Anthropic coverage remains via Twitter KOL, GitHub releases, and Reddit sources.\n- **Remove Third-Party RSS Sources section** from SKILL.md security docs (no longer applicable)\n\n## v3.3.0\n\n- **RSS Domain Validation**: New `expected_domains` field in sources.json rejects articles from unexpected origins (applied to anthropic-rss mirror)\n- **Email Shell Safety**: HTML body written to temp file before CLI delivery; subjects restricted to static format strings\n- **Discord Embed Suppression**: Footer links wrapped in `<>` to prevent preview embeds\n\n## v3.2.1\n\n- **Mandatory Reddit Execution**: Agent explicitly required to run `fetch-reddit.py` script — cannot skip or generate fake output\n\n## v3.2.0\n\n- **Unified English Templates**: All prompt instructions, section titles, stats footer, and example content standardized to English. Output language controlled by `<LANGUAGE>` placeholder at runtime.\n\n## v3.1.0\n\n- **Executive Summary**: 2-4 sentence overview of top stories at the beginning of each digest\n- **Community Buzz Section**: Merged Twitter/X Trending and Reddit Hot Discussions into unified 🔥 社区热议\n- **Reddit in Topic Sections**: Reddit posts now selected by quality_score alongside other sources\n- **Digest Footer Branding**: Shows skill version and OpenClaw link\n- **Prompt Fix**: Agent explicitly instructed to read Reddit data from merged JSON\n\n## v3.0.0\n\n- **Reddit Data Source**: New `fetch-reddit.py` script — 5th data layer using Reddit's public JSON API (no auth required). 13 subreddits: r/MachineLearning, r/LocalLLaMA, r/CryptoCurrency, r/artificial, r/ethereum, r/ChatGPT, r/singularity, r/OpenAI, r/Bitcoin, r/programming, r/Anthropic, r/defi, r/ExperiencedDevs\n- **Reddit Score Bonus**: Posts with score > 500 get +5, > 200 get +3, > 100 get +1 in quality scoring\n- **10 New Non-Reddit Sources**: Ben's Bites, The Decoder, a16z Crypto, Bankless (RSS); @ClementDelangue, @GregBrockman, @zuck (Twitter); MCP Servers, DeepSeek-V3, Meta Llama (GitHub)\n- **Tweet Engagement Metrics**: KOL entries display `👁|💬|🔁|❤️` stats in inline code blocks across all templates\n- **Date Timezone Fix**: Report date explicitly provided via `<DATE>` placeholder, preventing UTC/local mismatch\n- **Mandatory Links**: KOL Updates and Twitter/X Trending sections require source URLs for every entry\n- **Graceful Twitter Degradation**: Missing `X_BEARER_TOKEN` outputs empty JSON instead of failing\n- **URL Sanitization**: `resolve_link()` rejects non-HTTP(S) schemes\n- **Security Documentation**: Added Security Considerations section to SKILL.md\n- **Total Sources**: 132 (50 RSS + 47 Twitter + 22 GitHub + 13 Reddit + 4 web search topics)\n\n## v2.8.1\n\n- **Metrics Data Fix**: Agent now required to read actual `metrics` values from Twitter JSON data instead of defaulting to 0\n- **Email Template Enhancement**: Added KOL metrics and Twitter/X Trending section to email template\n\n## v2.8.0\n\n- **Tweet Metrics Display**: KOL entries show `👁|💬|🔁|❤️` engagement stats wrapped in inline code to prevent emoji enlargement on Discord\n- **Standardized Metrics Format**: Fixed 4-metric order, show 0 for missing values, one tweet per bullet with own URL\n- **10 New Sources (119 total)**: Ben's Bites, The Decoder, a16z Crypto, Bankless (RSS); @ClementDelangue, @GregBrockman, @zuck (Twitter); MCP Servers, DeepSeek-V3, Meta Llama (GitHub)\n\n## v2.7.0\n\n- **Tweet Engagement Metrics**: KOL Updates now display 👁 views, 💬 replies, 🔁 retweets, ❤️ likes from Twitter public_metrics across all templates (Discord, Email, Telegram)\n\n## v2.6.1\n\n- **Graceful Twitter Degradation**: Missing `X_BEARER_TOKEN` now outputs empty JSON and exits 0 instead of failing with exit code 1, allowing the pipeline to continue without Twitter data\n\n## v2.6.0\n\n- **Date Timezone Fix**: Added `<DATE>` placeholder to digest prompt — report date now explicitly provided by caller, preventing UTC/local timezone mismatch\n- **Mandatory Links in KOL/Trending**: KOL Updates and Twitter/X Trending sections now require source URLs for every entry (no link-free entries allowed)\n- **URL Sanitization**: `resolve_link()` in fetch-rss.py rejects non-HTTP(S) schemes (javascript:, data:, etc.)\n- **Third-Party Source Annotation**: Community-maintained RSS mirrors (e.g. anthropic-rss) are annotated with notes in sources.json\n- **Security Documentation**: Added Security Considerations section to SKILL.md covering shell execution model, input sanitization, and network access\n\n## v2.5.0\n\n- **Twitter Reply Filter Fix**: Use `referenced_tweets` field instead of text prefix to distinguish replies from mentions\n- **Scoring Consistency**: digest-prompt.md now matches code (`PENALTY_OLD_REPORT = -5`)\n- **Template Version Cleanup**: Removed hardcoded version numbers from email/markdown/telegram templates\n- **Article Count Fix**: `merge-sources.py` uses deduplicated count instead of inflated topic-grouped sum\n- **Pipeline Resume Support**: All fetch scripts support `--force` flag; skip if cached output < 1 hour old\n- **Source Health Monitoring**: New `scripts/source-health.py` tracks per-source success/failure history\n- **End-to-End Test**: New `scripts/test-pipeline.sh` smoke test for the full pipeline\n- **Archive Auto-Cleanup**: digest-prompt.md documents 90-day archive retention policy\n- **Twitter Rate Limiting**: Moved sleep into `fetch_user_tweets` for actual per-request rate limiting\n- **Web Article Scoring**: Web articles now use `calculate_base_score` instead of hardcoded 1.0\n- **Dead Code Removal**: Removed unused `load_sources_with_overlay` / `load_topics_with_overlay` wrappers\n\n## v2.4.0\n\n- **Batch Twitter Lookup**: Single API call for all username→ID resolution + 7-day local cache (~88→~45 API calls)\n- **Smart Dedup**: Token-based bucketing replaces O(n²) SequenceMatcher — only compares articles sharing 2+ key tokens\n- **Conditional Fetch (RSS)**: ETag/Last-Modified caching, 304 responses skip parsing\n- **Conditional Fetch (GitHub)**: Same caching pattern + prominent warning when GITHUB_TOKEN is unset\n- **`--no-cache` flag**: All fetch scripts support bypassing cache\n\n## v2.3.0\n\n- **GitHub Releases**: 19 tracked repositories as a fourth data source\n- **Data Source Stats Footer**: Pipeline statistics in all templates\n- **Twitter Queries**: Added to all 4 topics for better coverage\n- **Simplified Cron Prompts**: Reference digest-prompt.md with parameters only\n\n## v2.1.0\n\n- **Unified Source Model**: Single `sources.json` for RSS, Twitter, and web sources\n- **Enhanced Topics**: Richer topic definitions with search queries and filters\n- **Pipeline Scripts**: Modular fetch → merge → template workflow\n- **Quality Scoring**: Multi-source detection, deduplication, priority weighting\n- **Multiple Templates**: Discord, email, and markdown output formats\n- **Configuration Validation**: JSON schema validation and consistency checks\n- **User Customization**: Workspace config overrides for personalization\n\nFile v3.11.0:CONTRIBUTING.md\n\n# Contributing / Development Conventions\n\n## Version Management\n\n- **SemVer**: `SKILL.md` frontmatter `version` field is the single source of truth\n- **CHANGELOG.md**: reverse-chronological, update with every version bump\n- Every change must update **both** `SKILL.md version` + `CHANGELOG.md` + git commit & push\n- Changelog version format: `## v3.5.0` (prefixed with `v`)\n\n## Code Conventions\n\n- All prompts, templates, comments, and code in **English**\n- Output language controlled at runtime via `LANGUAGE` variable\n- Python: use `except Exception:` — never bare `except:`\n- No hardcoded credentials — all secrets via environment variables\n- When adding data sources, update `sources.json` schema **and** README source count\n\n## Security\n\n- ClawHub audit compliance: declare all `tools`/`bins`, file read/write paths, credential access in SKILL.md metadata\n- No third-party untrusted RSS mirrors (supply chain risk)\n- HTML email bodies written to temp files before CLI delivery\n- Subjects restricted to static format strings (no injection)\n- Discord embed suppression: wrap links in `<>` to prevent previews\n\n## Debugging\n\n- Full pipeline: `python3 scripts/run-pipeline.py --verbose --force`\n- Each step generates `*.meta.json` with timing, counts, and status\n- Individual scripts can be run standalone for targeted debugging\n\n## File Structure\n\n```\nSKILL.md          — Skill metadata (version, env vars, tools, files)\nCHANGELOG.md      — Version history\nREADME.md         — English docs\nREADME_CN.md      — Chinese docs\nconfig/defaults/  — Default sources.json, topics.json\nreferences/       — digest-prompt.md, output templates\nscripts/          — Python pipeline scripts\n```\n\n## Environment vs Code\n\n- **Never push environment-specific config to repo** — email sender names, API keys, file paths, channel IDs, timezone settings, etc. belong in local workspace config or env vars, not in skill code\n- Repo code uses `<PLACEHOLDER>` patterns; actual values are substituted at runtime\n- Local overrides go in `workspace/config/`, not in `config/defaults/`\n\n## Git Workflow\n\n- Commit messages: concise English, describe what changed\n- Push to `main` branch on github.com/draco-agent/tech-news-digest\n- No feature branches for solo development (direct to main)\n\nFile v3.11.0:README_CN.md\n\n# Tech News Digest\n\n> 自动化科技资讯汇总 — 138 个数据源，5 层管道，一句话安装。\n\n[English](README.md) | **中文**\n\n[![Tests](https://github.com/draco-agent/tech-news-digest/actions/workflows/test.yml/badge.svg)](https://github.com/draco-agent/tech-news-digest/actions/workflows/test.yml)\n[![Python 3.8+](https://img.shields.io/badge/python-3.8+-blue.svg)](https://www.python.org/downloads/)\n[![ClawHub](https://img.shields.io/badge/ClawHub-tech--news--digest-blueviolet)](https://clawhub.com/draco-agent/tech-news-digest)\n[![MIT License](https://img.shields.io/badge/license-MIT-green.svg)](LICENSE)\n\n## 💬 一句话安装\n\n跟你的 [OpenClaw](https://openclaw.ai) AI 助手说：\n\n> **\"安装 tech-news-digest，每天早上 9 点发科技日报到 #tech-news 频道\"**\n\n搞定。Bot 会自动安装、配置、定时、推送——全程对话完成。\n\n更多示例：\n\n> 🗣️ \"配置一个每周 AI 周报，只要 LLM 和 AI Agent 板块，每周一发到 Discord #ai-weekly\"\n\n> 🗣️ \"安装 tech-news-digest，加上我的 RSS 源，发送科技新闻到 Telegram\"\n\n> 🗣️ \"现在就给我生成一份科技日报，跳过 Twitter 数据源\"\n\n或通过 CLI 安装：\n```bash\nclawhub install tech-news-digest\n```\n\n## 📊 你会得到什么\n\n基于 **138 个数据源** 的质量评分、去重科技日报：\n\n| 层级 | 数量 | 内容 |\n|------|------|------|\n| 📡 RSS | 49 个订阅源 | OpenAI、Anthropic、Ben's Bites、HN、36氪、CoinDesk… |\n| 🐦 Twitter/X | 48 个 KOL | @karpathy、@VitalikButerin、@sama、@elonmusk… |\n| 🔍 Web 搜索 | 4 个主题 | Brave Search API + 时效过滤 |\n| 🐙 GitHub | 28 个仓库 | 关键项目的 Release 跟踪（LangChain、vLLM、DeepSeek、Llama…） |\n| 🗣️ Reddit | 13 个子版块 | r/MachineLearning、r/LocalLLaMA、r/CryptoCurrency… |\n\n### 数据管道\n\n```\n       run-pipeline.py (~30秒)\n              ↓\n  RSS ─┐\n  Twitter ─┤\n  Web ─────┤── 并行采集 ──→ merge-sources.py\n  GitHub ──┤\n  Reddit ──┘\n              ↓\n  质量评分 → 去重 → 主题分组\n              ↓\n    Discord / 邮件 / PDF 输出\n```\n\n**质量评分**：优先级源 (+3)、多源交叉验证 (+5)、时效性 (+2)、互动度 (+1~+5)、Reddit 热度加分 (+1/+3/+5)、已报道过 (-5)。\n\n## ⚙️ 配置\n\n- `config/defaults/sources.json` — 138 个内置数据源\n- `config/defaults/topics.json` — 4 个主题，含搜索查询和 Twitter 查询\n- 用户自定义配置放 `workspace/config/`，优先级更高\n\n## 🎨 自定义数据源\n\n开箱即用，内置 138 个数据源——但完全可自定义。将默认配置复制到 workspace 并覆盖：\n\n```bash\n# 复制并自定义\ncp config/defaults/sources.json workspace/config/tech-news-digest-sources.json\ncp config/defaults/topics.json workspace/config/tech-news-digest-topics.json\n```\n\n你的配置文件会与默认配置**合并**：\n- **覆盖**：`id` 匹配的源会被你的版本替换\n- **新增**：使用新的 `id` 即可添加自定义源\n- **禁用**：对匹配的 `id` 设置 `\"enabled\": false`\n\n```json\n{\n  \"sources\": [\n    {\"id\": \"my-blog\", \"type\": \"rss\", \"enabled\": true, \"url\": \"https://myblog.com/feed\", \"topics\": [\"llm\"]},\n    {\"id\": \"openai-blog\", \"enabled\": false}\n  ]\n}\n```\n\n不需要复制整个文件——只写你要改的部分。\n\n## 🔧 可选配置\n\n所有环境变量均为可选，管道会自动使用可用的数据源。\n\n```bash\nexport TWITTERAPI_IO_KEY=\"...\"    # twitterapi.io (~$5/月) — 启用 Twitter 数据层\nexport X_BEARER_TOKEN=\"...\"       # Twitter/X 官方 API — 备选 Twitter 后端\nexport TWITTER_API_BACKEND=\"auto\" # auto|twitterapiio|official（默认: auto）\nexport TAVILY_API_KEY=\"tvly-xxx\"  # Tavily Search API（替代方案，免费 1000 次/月）\nexport BRAVE_API_KEYS=\"k1,k2,k3\" # Brave Search API 密钥（逗号分隔，自动轮换）\nexport BRAVE_API_KEY=\"...\"        # 单密钥回退\nexport BRAVE_PLAN=\"free\"          # 覆盖速率限制检测: free|pro\nexport WEB_SEARCH_BACKEND=\"auto\" # auto|brave|tavily（默认: auto）\nexport GITHUB_TOKEN=\"...\"         # GitHub API — 提高速率限制（未设置时自动从 GitHub App 生成）\npip install weasyprint             # 启用 PDF 报告生成\n```\n\n## 🧪 测试\n\n```bash\npython -m unittest discover -s tests -v   # 41 个测试，纯标准库\n```\n\n## 📂 仓库地址\n\n**GitHub**: [github.com/draco-agent/tech-news-digest](https://github.com/draco-agent/tech-news-digest)\n\n## 📄 开源协议\n\nMIT License — 详见 [LICENSE](LICENSE)\n\nFile v3.11.0:config/defaults/sources.json\n\n{\n  \"_description\": \"Unified data sources configuration. RSS feeds, Twitter/X KOLs, and web search sources. Each source binds to topics and has enabled/priority fields.\",\n  \"_updated\": \"2025-06-02\",\n  \"_version\": \"2.5.0\",\n  \"sources\": [\n    {\n      \"id\": \"simonwillison-rss\",\n      \"type\": \"rss\",\n      \"name\": \"Simon Willison\",\n      \"url\": \"https://simonwillison.net/atom/everything/\",\n      \"enabled\": true,\n      \"priority\": true,\n      \"topics\": [\n        \"llm\",\n        \"ai-agent\",\n        \"frontier-tech\"\n      ],\n      \"note\": \"LLM/AI tooling, prolific blogger\"\n    },\n    {\n      \"id\": \"garymarcus-rss\",\n      \"type\": \"rss\",\n      \"name\": \"Gary Marcus\",\n      \"url\": \"https://garymarcus.substack.com/feed\",\n      \"enabled\": true,\n      \"priority\": true,\n      \"topics\": [\n        \"llm\",\n        \"frontier-tech\"\n      ],\n      \"note\": \"AI critic, industry analysis\"\n    },\n    {\n      \"id\": \"huggingface-rss\",\n      \"type\": \"rss\",\n      \"name\": \"Hugging Face Blog\",\n      \"url\": \"https://huggingface.co/blog/feed.xml\",\n      \"enabled\": true,\n      \"priority\": true,\n      \"topics\": [\n        \"llm\",\n        \"ai-agent\"\n      ],\n      \"note\": \"Open source AI/ML\"\n    },\n    {\n      \"id\": \"openai-rss\",\n      \"type\": \"rss\",\n      \"name\": \"OpenAI Blog\",\n      \"url\": \"https://openai.com/blog/rss.xml\",\n      \"enabled\": true,\n      \"priority\": true,\n      \"topics\": [\n        \"llm\",\n        \"ai-agent\",\n        \"frontier-tech\"\n      ],\n      \"note\": \"Official OpenAI updates\"\n    },\n    {\n      \"id\": \"sebas-rss\",\n      \"type\": \"rss\",\n      \"name\": \"Sebastian Raschka\",\n      \"url\": \"https://magazine.sebastianraschka.com/feed\",\n      \"enabled\": true,\n      \"priority\": true,\n      \"topics\": [\n        \"llm\",\n        \"frontier-tech\"\n      ],\n      \"note\": \"LLM research and tutorials\"\n    },\n    {\n      \"id\": \"lilian-rss\",\n      \"type\": \"rss\",\n      \"name\": \"Lil'Log (Lilian Weng)\",\n      \"url\": \"https://lilianweng.github.io/index.xml\",\n      \"enabled\": true,\n      \"priority\": false,\n      \"topics\": [\n        \"llm\",\n        \"frontier-tech\"\n      ],\n      \"note\": \"In-depth ML tutorials\"\n    },\n    {\n      \"id\": \"gwern-rss\",\n      \"type\": \"rss\",\n      \"name\": \"Gwern\",\n      \"url\": \"https://gwern.substack.com/feed\",\n      \"enabled\": true,\n      \"priority\": false,\n      \"topics\": [\n        \"llm\",\n        \"frontier-tech\"\n      ],\n      \"note\": \"Deep AI/ML research essays\"\n    },\n    {\n      \"id\": \"dwarkesh-rss\",\n      \"type\": \"rss\",\n      \"name\": \"Dwarkesh Patel\",\n      \"url\": \"https://www.dwarkeshpatel.com/feed\",\n      \"enabled\": true,\n      \"priority\": false,\n      \"topics\": [\n        \"llm\",\n        \"ai-agent\"\n      ],\n      \"note\": \"AI interviews and analysis\"\n    },\n    {\n      \"id\": \"minimaxir-rss\",\n      \"type\": \"rss\",\n      \"name\": \"minimaxir (Max Woolf)\",\n      \"url\": \"https://minimaxir.com/index.xml\",\n      \"enabled\": true,\n      \"priority\": false,\n      \"topics\": [\n        \"llm\",\n        \"frontier-tech\"\n      ],\n      \"note\": \"LLM benchmarks and experiments\"\n    },\n    {\n      \"id\": \"googleai-rss\",\n      \"type\": \"rss\",\n      \"name\": \"Google AI Blog\",\n      \"url\": \"https://blog.google/technology/ai/rss/\",\n      \"enabled\": true,\n      \"priority\": false,\n      \"topics\": [\n        \"llm\",\n        \"ai-agent\",\n        \"frontier-tech\"\n      ],\n      \"note\": \"Google AI research\"\n    },\n    {\n      \"id\": \"vitalik-rss\",\n      \"type\": \"rss\",\n      \"name\": \"Vitalik Buterin\",\n      \"url\": \"https://vitalik.eth.limo/feed.xml\",\n      \"enabled\": true,\n      \"priority\": true,\n      \"topics\": [\n        \"crypto\"\n      ],\n      \"note\": \"Ethereum founder\"\n    },\n    {\n      \"id\": \"coindesk-rss\",\n      \"type\": \"rss\",\n      \"name\": \"CoinDesk\",\n      \"url\": \"https://www.coindesk.com/arc/outboundfeeds/rss/\",\n      \"enabled\": true,\n      \"priority\": true,\n      \"topics\": [\n        \"crypto\"\n      ],\n      \"note\": \"Major crypto news\"\n    },\n    {\n      \"id\": \"theblock-rss\",\n      \"type\": \"rss\",\n      \"name\": \"The Block\",\n      \"url\": \"https://www.theblock.co/rss.xml\",\n      \"enabled\": true,\n      \"priority\": true,\n      \"topics\": [\n        \"crypto\"\n      ],\n      \"note\": \"Crypto news\"\n    },\n    {\n      \"id\": \"decrypt-rss\",\n      \"type\": \"rss\",\n      \"name\": \"Decrypt\",\n      \"url\": \"https://decrypt.co/feed\",\n      \"enabled\": true,\n      \"priority\": false,\n      \"topics\": [\n        \"crypto\"\n      ],\n      \"note\": \"Crypto and web3 news\"\n    },\n    {\n      \"id\": \"cointelegraph-rss\",\n      \"type\": \"rss\",\n      \"name\": \"Cointelegraph\",\n      \"url\": \"https://cointelegraph.com/rss\",\n      \"enabled\": true,\n      \"priority\": false,\n      \"topics\": [\n        \"crypto\"\n      ],\n      \"note\": \"Crypto news and analysis\"\n    },\n    {\n      \"id\": \"hn-rss\",\n      \"type\": \"rss\",\n      \"name\": \"Hacker News Frontpage\",\n      \"url\": \"https://hnrss.org/frontpage\",\n      \"enabled\": true,\n      \"priority\": true,\n      \"topics\": [\n        \"frontier-tech\"\n      ],\n      \"note\": \"HN top stories\"\n    },\n    {\n      \"id\": \"ars-rss\",\n      \"type\": \"rss\",\n      \"name\": \"Ars Technica\",\n      \"url\": \"https://feeds.arstechnica.com/arstechnica/index\",\n      \"enabled\": true,\n      \"priority\": false,\n      \"topics\": [\n        \"frontier-tech\"\n      ],\n      \"note\": \"Tech news\"\n    },\n    {\n      \"id\": \"techcrunch-rss\",\n      \"type\": \"rss\",\n      \"name\": \"TechCrunch\",\n      \"url\": \"https://techcrunch.com/feed/\",\n      \"enabled\": true,\n      \"priority\": false,\n      \"topics\": [\n        \"frontier-tech\"\n      ],\n      \"note\": \"Startup and tech news\"\n    },\n    {\n      \"id\": \"verge-rss\",\n      \"type\": \"rss\",\n      \"name\": \"The Verge\",\n      \"url\": \"https://www.theverge.com/rss/index.xml\",\n      \"enabled\": true,\n      \"priority\": false,\n      \"topics\": [\n        \"frontier-tech\"\n      ],\n      \"note\": \"Tech news\"\n    },\n    {\n      \"id\": \"krebs-rss\",\n      \"type\": \"rss\",\n      \"name\": \"Krebs on Security\",\n      \"url\": \"https://krebsonsecurity.com/feed/\",\n      \"enabled\": true,\n      \"priority\": false,\n      \"topics\": [\n        \"frontier-tech\"\n      ],\n      \"note\": \"Cybersecurity\"\n    },\n    {\n      \"id\": \"daringfireball-rss\",\n      \"type\": \"rss\",\n      \"name\": \"Daring Fireball\",\n      \"url\": \"https://daringfireball.net/feeds/main\",\n      \"enabled\": true,\n      \"priority\": false,\n      \"topics\": [\n        \"frontier-tech\"\n      ],\n      \"note\": \"Apple/tech commentary\"\n    },\n    {\n      \"id\": \"pg-rss\",\n      \"type\": \"rss\",\n      \"name\": \"Paul Graham\",\n      \"url\": \"http://www.aaronsw.com/2002/feeds/pgessays.rss\",\n      \"enabled\": true,\n      \"priority\": false,\n      \"topics\": [\n        \"frontier-tech\"\n      ],\n      \"note\": \"Startups and tech essays\"\n    },\n    {\n      \"id\": \"troyhunt-rss\",\n      \"type\": \"rss\",\n      \"name\": \"Troy Hunt\",\n      \"url\": \"https://www.troyhunt.com/rss/\",\n      \"enabled\": true,\n      \"priority\": false,\n      \"topics\": [\n        \"frontier-tech\"\n      ],\n      \"note\": \"Security, HIBP\"\n    },\n    {\n      \"id\": \"antirez-rss\",\n      \"type\": \"rss\",\n      \"name\": \"antirez\",\n      \"url\": \"http://antirez.com/rss\",\n      \"enabled\": true,\n      \"priority\": false,\n      \"topics\": [\n        \"frontier-tech\"\n      ],\n      \"note\": \"Redis creator, systems\"\n    },\n    {\n      \"id\": \"mitchellh-rss\",\n      \"type\": \"rss\",\n      \"name\": \"Mitchell Hashimoto\",\n      \"url\": \"https://mitchellh.com/feed.xml\",\n      \"enabled\": true,\n      \"priority\": false,\n      \"topics\": [\n        \"frontier-tech\"\n      ],\n      \"note\": \"Ghostty, infrastructure\"\n    },\n    {\n      \"id\": \"geohot-rss\",\n      \"type\": \"rss\",\n      \"name\": \"Geohot\",\n      \"url\": \"https://geohot.github.io/blog/feed.xml\",\n      \"enabled\": true,\n      \"priority\": false,\n      \"topics\": [\n        \"frontier-tech\",\n        \"llm\"\n      ],\n      \"note\": \"tinygrad, AI infrastructure\"\n    },\n    {\n      \"id\": \"ml-reddit-rss\",\n      \"type\": \"rss\",\n      \"name\": \"r/MachineLearning\",\n      \"url\": \"https://www.reddit.com/r/MachineLearning/.rss\",\n      \"enabled\": true,\n      \"priority\": false,\n      \"topics\": [\n        \"llm\",\n        \"frontier-tech\"\n      ],\n      \"note\": \"Reddit ML community\"\n    },\n    {\n      \"id\": \"36kr-rss\",\n      \"type\": \"rss\",\n      \"name\": \"36氪\",\n      \"url\": \"https://36kr.com/feed\",\n      \"enabled\": true,\n      \"priority\": true,\n      \"topics\": [\n        \"frontier-tech\",\n        \"crypto\"\n      ],\n      \"note\": \"中文科技媒体\"\n    },\n    {\n      \"id\": \"synced-rss\",\n      \"type\": \"rss\",\n      \"name\": \"机器之心 Synced\",\n      \"url\": \"https://www.jiqizhixin.com/rss\",\n      \"enabled\": true,\n      \"priority\": true,\n      \"topics\": [\n        \"llm\",\n        \"ai-agent\",\n        \"frontier-tech\"\n      ],\n      \"note\": \"中文AI媒体\"\n    },\n    {\n      \"id\": \"qbitai-rss\",\n      \"type\": \"rss\",\n      \"name\": \"量子位 QbitAI\",\n      \"url\": \"https://www.qbitai.com/feed\",\n      \"enabled\": true,\n      \"priority\": false,\n      \"topics\": [\n        \"llm\",\n        \"ai-agent\"\n      ],\n      \"note\": \"中文AI媒体 (may 403)\"\n    },\n    {\n      \"id\": \"infoq-rss\",\n      \"type\": \"rss\",\n      \"name\": \"InfoQ 中文\",\n      \"url\": \"https://www.infoq.cn/feed\",\n      \"enabled\": true,\n      \"priority\": false,\n      \"topics\": [\n        \"frontier-tech\"\n      ],\n      \"note\": \"技术社区\"\n    },\n    {\n      \"id\": \"sama-twitter\",\n      \"type\": \"twitter\",\n      \"name\": \"Sam Altman (OpenAI CEO)\",\n      \"handle\": \"sama\",\n      \"enabled\": true,\n      \"priority\": true,\n      \"topics\": [\n        \"llm\",\n        \"ai-agent\",\n        \"frontier-tech\"\n      ],\n      \"note\": \"OpenAI CEO\"\n    },\n    {\n      \"id\": \"openai-twitter\",\n      \"type\": \"twitter\",\n      \"name\": \"OpenAI official\",\n      \"handle\": \"OpenAI\",\n      \"enabled\": true,\n      \"priority\": true,\n      \"topics\": [\n        \"llm\",\n        \"ai-agent\",\n        \"frontier-tech\"\n      ],\n      \"note\": \"OpenAI official\"\n    },\n    {\n      \"id\": \"anthropic-twitter\",\n      \"type\": \"twitter\",\n      \"name\": \"Anthropic official\",\n      \"handle\": \"AnthropicAI\",\n      \"enabled\": true,\n      \"priority\": true,\n      \"topics\": [\n        \"llm\",\n        \"ai-agent\",\n        \"frontier-tech\"\n      ],\n      \"note\": \"Anthropic official\"\n    },\n    {\n      \"id\": \"ylecun-twitter\",\n      \"type\": \"twitter\",\n      \"name\": \"Yann LeCun (Meta AI)\",\n      \"handle\": \"ylecun\",\n      \"enabled\": true,\n      \"priority\": true,\n      \"topics\": [\n        \"llm\",\n        \"frontier-tech\"\n      ],\n      \"note\": \"Meta AI\"\n    },\n    {\n      \"id\": \"mistral-twitter\",\n      \"type\": \"twitter\",\n      \"name\": \"Mistral AI official\",\n      \"handle\": \"MistralAI\",\n      \"enabled\": true,\n      \"priority\": true,\n      \"topics\": [\n        \"llm\",\n        \"frontier-tech\"\n      ],\n      \"note\": \"Mistral AI official\"\n    },\n    {\n      \"id\": \"deepmind-twitter\",\n      \"type\": \"twitter\",\n      \"name\": \"Google DeepMind official\",\n      \"handle\": \"GoogleDeepMind\",\n      \"enabled\": true,\n      \"priority\": true,\n      \"topics\": [\n        \"llm\",\n        \"frontier-tech\"\n      ],\n      \"note\": \"Google DeepMind official\"\n    },\n    {\n      \"id\": \"googleai-twitter\",\n      \"type\": \"twitter\",\n      \"name\": \"Google AI official\",\n      \"handle\": \"GoogleAI\",\n      \"enabled\": true,\n      \"priority\": true,\n      \"topics\": [\n        \"llm\",\n        \"ai-agent\",\n        \"frontier-tech\"\n      ],\n      \"note\": \"Google AI official\"\n    },\n    {\n      \"id\": \"xai-twitter\",\n      \"type\": \"twitter\",\n      \"name\": \"xAI official\",\n      \"handle\": \"xai\",\n      \"enabled\": true,\n      \"priority\": true,\n      \"topics\": [\n        \"llm\",\n        \"ai-agent\",\n        \"frontier-tech\"\n      ],\n      \"note\": \"xAI official\"\n    },\n    {\n      \"id\": \"karpathy-twitter\",\n      \"type\": \"twitter\",\n      \"name\": \"Andrej Karpathy\",\n      \"handle\": \"karpathy\",\n      \"enabled\": true,\n      \"priority\": true,\n      \"topics\": [\n        \"llm\",\n        \"ai-agent\",\n        \"frontier-tech\"\n      ],\n      \"note\": \"AI researcher\"\n    },\n    {\n      \"id\": \"andrewng-twitter\",\n      \"type\": \"twitter\",\n      \"name\": \"Andrew Ng\",\n      \"handle\": \"AndrewYNg\",\n      \"enabled\": true,\n      \"priority\": true,\n      \"topics\": [\n        \"llm\",\n        \"ai-agent\"\n      ],\n      \"note\": \"AI educator\"\n    },\n    {\n      \"id\": \"jimfan-twitter\",\n      \"type\": \"twitter\",\n      \"name\": \"Jim Fan (NVIDIA)\",\n      \"handle\": \"DrJimFan\",\n      \"enabled\": true,\n      \"priority\": true,\n      \"topics\": [\n        \"ai-agent\",\n        \"frontier-tech\"\n      ],\n      \"note\": \"NVIDIA AI\"\n    },\n    {\n      \"id\": \"hf-twitter\",\n      \"type\": \"twitter\",\n      \"name\": \"Hugging Face official\",\n      \"handle\": \"huggingface\",\n      \"enabled\": true,\n      \"priority\": true,\n      \"topics\": [\n        \"llm\",\n        \"ai-agent\"\n      ],\n      \"note\": \"Hugging Face official\"\n    },\n    {\n      \"id\": \"langchain-twitter\",\n      \"type\": \"twitter\",\n      \"name\": \"LangChain official\",\n      \"handle\": \"LangChain\",\n      \"enabled\": true,\n      \"priority\": false,\n      \"topics\": [\n        \"ai-agent\"\n      ],\n      \"note\": \"LangChain official\"\n    },\n    {\n      \"id\": \"llamaindex-twitter\",\n      \"type\": \"twitter\",\n      \"name\": \"LlamaIndex official\",\n      \"handle\": \"llama_index\",\n      \"enabled\": true,\n      \"priority\": false,\n      \"topics\": [\n        \"ai-agent\"\n      ],\n      \"note\": \"LlamaIndex official\"\n    },\n    {\n      \"id\": \"emad-twitter\",\n      \"type\": \"twitter\",\n      \"name\": \"Emad Mostaque\",\n      \"handle\": \"EMostaque\",\n      \"enabled\": true,\n      \"priority\": false,\n      \"topics\": [\n        \"llm\",\n        \"frontier-tech\"\n      ],\n      \"note\": \"Stability AI\"\n    },\n    {\n      \"id\": \"sebastian-twitter\",\n      \"type\": \"twitter\",\n      \"name\": \"Sebastian Raschka\",\n      \"handle\": \"rasbt\",\n      \"enabled\": true,\n      \"priority\": false,\n      \"topics\": [\n        \"llm\",\n        \"frontier-tech\"\n      ],\n      \"note\": \"AI researcher\"\n    },\n    {\n      \"id\": \"vitalik-twitter\",\n      \"type\": \"twitter\",\n      \"name\": \"Vitalik Buterin (Ethereum)\",\n      \"handle\": \"VitalikButerin\",\n      \"enabled\": true,\n      \"priority\": true,\n      \"topics\": [\n        \"crypto\"\n      ],\n      \"note\": \"Ethereum founder\"\n    },\n    {\n      \"id\": \"cz-twitter\",\n      \"type\": \"twitter\",\n      \"name\": \"CZ (Binance)\",\n      \"handle\": \"cz_binance\",\n      \"enabled\": true,\n      \"priority\": true,\n      \"topics\": [\n        \"crypto\"\n      ],\n      \"note\": \"Binance\"\n    },\n    {\n      \"id\": \"brian-twitter\",\n      \"type\": \"twitter\",\n      \"name\": \"Brian Armstrong (Coinbase)\",\n      \"handle\": \"brian_armstrong\",\n      \"enabled\": true,\n      \"priority\": true,\n      \"topics\": [\n        \"crypto\"\n      ],\n      \"note\": \"Coinbase\"\n    },\n    {\n      \"id\": \"saylor-twitter\",\n      \"type\": \"twitter\",\n      \"name\": \"Michael Saylor (MicroStrategy)\",\n      \"handle\": \"saylor\",\n      \"enabled\": true,\n      \"priority\": true,\n      \"topics\": [\n        \"crypto\"\n      ],\n      \"note\": \"MicroStrategy\"\n    },\n    {\n      \"id\": \"pomp-twitter\",\n      \"type\": \"twitter\",\n      \"name\": \"Anthony Pompliano\",\n      \"handle\": \"APompliano\",\n      \"enabled\": true,\n      \"priority\": false,\n      \"topics\": [\n        \"crypto\"\n      ],\n      \"note\": \"Crypto influencer\"\n    },\n    {\n      \"id\": \"zachxbt-twitter\",\n      \"type\": \"twitter\",\n      \"name\": \"ZachXBT\",\n      \"handle\": \"zachxbt\",\n      \"enabled\": true,\n      \"priority\": false,\n      \"topics\": [\n        \"crypto\"\n      ],\n      \"note\": \"on-chain investigator\"\n    },\n    {\n      \"id\": \"wu-twitter\",\n      \"type\": \"twitter\",\n      \"name\": \"Wu Blockchain\",\n      \"handle\": \"WuBlockchain\",\n      \"enabled\": true,\n      \"priority\": true,\n      \"topics\": [\n        \"crypto\"\n      ],\n      \"note\": \"吴说区块链\"\n    },\n    {\n      \"id\": \"discus-twitter\",\n      \"type\": \"twitter\",\n      \"name\": \"神鱼 DiscusFish\",\n      \"handle\": \"bitfish\",\n      \"enabled\": true,\n      \"priority\": false,\n      \"topics\": [\n        \"crypto\"\n      ],\n      \"note\": \"F2Pool/Cobo co-founder\"\n    },\n    {\n      \"id\": \"mindao-twitter\",\n      \"type\": \"twitter\",\n      \"name\": \"Mindao\",\n      \"handle\": \"mindaoyang\",\n      \"enabled\": true,\n      \"priority\": false,\n      \"topics\": [\n        \"crypto\"\n      ],\n      \"note\": \"dForce founder\"\n    },\n    {\n      \"id\": \"herbert-twitter\",\n      \"type\": \"twitter\",\n      \"name\": \"Herbert\",\n      \"handle\": \"herbertcrypto\",\n      \"enabled\": true,\n      \"priority\": false,\n      \"topics\": [\n        \"crypto\"\n      ],\n      \"note\": \"PANews founder\"\n    },\n    {\n      \"id\": \"elon-twitter\",\n      \"type\": \"twitter\",\n      \"name\": \"Elon Musk\",\n      \"handle\": \"elonmusk\",\n      \"enabled\": true,\n      \"priority\": false,\n      \"topics\": [\n        \"frontier-tech\"\n      ],\n      \"note\": \"Entrepreneur\"\n    },\n    {\n      \"id\": \"sundar-twitter\",\n      \"type\": \"twitter\",\n      \"name\": \"Sundar Pichai\",\n      \"handle\": \"sundarpichai\",\n      \"enabled\": true,\n      \"priority\": false,\n      \"topics\": [\n        \"frontier-tech\"\n      ],\n      \"note\": \"Google CEO\"\n    },\n    {\n      \"id\": \"pmarca-twitter\",\n      \"type\": \"twitter\",\n      \"name\": \"Marc Andreessen\",\n      \"handle\": \"pmarca\",\n      \"enabled\": true,\n      \"priority\": false,\n      \"topics\": [\n        \"frontier-tech\",\n        \"crypto\"\n      ],\n      \"note\": \"a16z\"\n    },\n    {\n      \"id\": \"levie-twitter\",\n      \"type\": \"twitter\",\n      \"name\": \"Aaron Levie\",\n      \"handle\": \"levie\",\n      \"enabled\": true,\n      \"priority\": false,\n      \"topics\": [\n        \"frontier-tech\"\n      ],\n      \"note\": \"Box CEO\"\n    },\n    {\n      \"id\": \"satya-twitter\",\n      \"type\": \"twitter\",\n      \"name\": \"Satya Nadella\",\n      \"handle\": \"satyanadella\",\n      \"enabled\": true,\n      \"priority\": false,\n      \"topics\": [\n        \"frontier-tech\"\n      ],\n      \"note\": \"Microsoft CEO\"\n    },\n    {\n      \"id\": \"mit-tech-review-rss\",\n      \"type\": \"rss\",\n      \"name\": \"MIT Technology Review\",\n      \"url\": \"https://www.technologyreview.com/feed\",\n      \"enabled\": true,\n      \"priority\": true,\n      \"topics\": [\n        \"llm\",\n        \"ai-agent\",\n        \"frontier-tech\"\n      ],\n      \"note\": \"AI policy + deep analysis\"\n    },\n    {\n      \"id\": \"venturebeat-ai-rss\",\n      \"type\": \"rss\",\n      \"name\": \"VentureBeat AI\",\n      \"url\": \"https://venturebeat.com/category/ai/feed/\",\n      \"enabled\": true,\n      \"priority\": false,\n      \"topics\": [\n        \"llm\",\n        \"ai-agent\"\n      ],\n      \"note\": \"AI industry news\"\n    },\n    {\n      \"id\": \"404media-rss\",\n      \"type\": \"rss\",\n      \"name\": \"404 Media\",\n      \"url\": \"https://www.404media.co/rss\",\n      \"enabled\": true,\n      \"priority\": false,\n      \"topics\": [\n        \"frontier-tech\"\n      ],\n      \"note\": \"Independent tech investigative journalism\"\n    },\n    {\n      \"id\": \"aisnakeoil-rss\",\n      \"type\": \"rss\",\n      \"name\": \"AI Snake Oil\",\n      \"url\": \"https://aisnakeoil.substack.com/feed\",\n      \"enabled\": true,\n      \"priority\": false,\n      \"topics\": [\n        \"llm\",\n        \"frontier-tech\"\n      ],\n      \"note\": \"Princeton professor, critical AI perspective\"\n    },\n    {\n      \"id\": \"bytebytego-rss\",\n      \"type\": \"rss\",\n      \"name\": \"ByteByteGo\",\n      \"url\": \"https://blog.bytebytego.com/feed\",\n      \"enabled\": true,\n      \"priority\": false,\n      \"topics\": [\n        \"frontier-tech\"\n      ],\n      \"note\": \"System design + engineering by Alex Xu\"\n    },\n    {\n      \"id\": \"nvidia-ai-rss\",\n      \"type\": \"rss\",\n      \"name\": \"NVIDIA AI Blog\",\n      \"url\": \"https://blogs.nvidia.com/feed/\",\n      \"enabled\": true,\n      \"priority\": true,\n      \"topics\": [\n        \"llm\",\n        \"frontier-tech\"\n      ],\n      \"note\": \"GPU/AI infrastructure\"\n    },\n    {\n      \"id\": \"deepmind-rss\",\n      \"type\": \"rss\",\n      \"name\": \"Google DeepMind Blog\",\n      \"url\": \"https://deepmind.google/blog/rss.xml\",\n      \"enabled\": true,\n      \"priority\": true,\n      \"topics\": [\n        \"llm\",\n        \"frontier-tech\"\n      ],\n      \"note\": \"Frontier AI research\"\n    },\n    {\n      \"id\": \"producthunt-rss\",\n      \"type\": \"rss\",\n      \"name\": \"Product Hunt\",\n      \"url\": \"https://www.producthunt.com/feed\",\n      \"enabled\": true,\n      \"priority\": false,\n      \"topics\": [\n        \"frontier-tech\"\n      ],\n      \"note\": \"New product discovery, AI tools\"\n    },\n    {\n      \"id\": \"messari-rss\",\n      \"type\": \"rss\",\n      \"name\": \"Messari\",\n      \"url\": \"https://messari.io/rss\",\n      \"enabled\": true,\n      \"priority\": false,\n      \"topics\": [\n        \"crypto\"\n      ],\n      \"note\": \"Crypto research reports\"\n    },\n    {\n      \"id\": \"defiant-rss\",\n      \"type\": \"rss\",\n      \"name\": \"The Defiant\",\n      \"url\": \"https://thedefiant.io/feed\",\n      \"enabled\": true,\n      \"priority\": false,\n      \"topics\": [\n        \"crypto\"\n      ],\n      \"note\": \"DeFi professional media\"\n    },\n    {\n      \"id\": \"ifanr-rss\",\n      \"type\": \"rss\",\n      \"name\": \"爱范儿\",\n      \"url\": \"https://www.ifanr.com/feed\",\n      \"enabled\": true,\n      \"priority\": false,\n      \"topics\": [\n        \"frontier-tech\"\n      ],\n      \"note\": \"Chinese tech product reviews\"\n    },\n    {\n      \"id\": \"sspai-rss\",\n      \"type\": \"rss\",\n      \"name\": \"少数派\",\n      \"url\": \"https://sspai.com/feed\",\n      \"enabled\": true,\n      \"priority\": false,\n      \"topics\": [\n        \"frontier-tech\"\n      ],\n      \"note\": \"Chinese productivity + tech depth\"\n    },\n    {\n      \"id\": \"wired-rss\",\n      \"type\": \"rss\",\n      \"name\": \"Wired\",\n      \"url\": \"https://www.wired.com/feed/rss\",\n      \"enabled\": true,\n      \"priority\": false,\n      \"topics\": [\n        \"frontier-tech\"\n      ],\n      \"note\": \"Tech culture deep reporting\"\n    },\n    {\n      \"id\": \"ieee-spectrum-rss\",\n      \"type\": \"rss\",\n      \"name\": \"IEEE Spectrum\",\n      \"url\": \"https://spectrum.ieee.org/feeds/feed.rss\",\n      \"enabled\": true,\n      \"priority\": false,\n      \"topics\": [\n        \"frontier-tech\"\n      ],\n      \"note\": \"Engineering + frontier tech authority\"\n    },\n    {\n      \"id\": \"rowancheung-twitter\",\n      \"type\": \"twitter\",\n      \"name\": \"Rowan Cheung (The Rundown AI)\",\n      \"handle\": \"rowancheung\",\n      \"enabled\": true,\n      \"priority\": true,\n      \"topics\": [\n        \"llm\",\n        \"ai-agent\"\n      ],\n      \"note\": \"AI newsletter founder, 567K followers\"\n    },\n    {\n      \"id\": \"yudkowsky-twitter\",\n      \"type\": \"twitter\",\n      \"name\": \"Eliezer Yudkowsky\",\n      \"handle\": \"ESYudkowsky\",\n      \"enabled\": true,\n      \"priority\": false,\n      \"topics\": [\n        \"llm\",\n        \"frontier-tech\"\n      ],\n      \"note\": \"AI safety pioneer\"\n    },\n    {\n      \"id\": \"demis-twitter\",\n      \"type\": \"twitter\",\n      \"name\": \"Demis Hassabis (DeepMind CEO)\",\n      \"handle\": \"demishassabis\",\n      \"enabled\": true,\n      \"priority\": true,\n      \"topics\": [\n        \"llm\",\n        \"frontier-tech\"\n      ],\n      \"note\": \"DeepMind CEO, Nobel laureate\"\n    },\n    {\n      \"id\": \"dario-twitter\",\n      \"type\": \"twitter\",\n      \"name\": \"Dario Amodei (Anthropic CEO)\",\n      \"handle\": \"DarioAmodei\",\n      \"enabled\": true,\n      \"priority\": true,\n      \"topics\": [\n        \"llm\",\n        \"ai-agent\"\n      ],\n      \"note\": \"Anthropic CEO\"\n    },\n    {\n      \"id\": \"hwchase-twitter\",\n      \"type\": \"twitter\",\n      \"name\": \"Harrison Chase (LangChain)\",\n      \"handle\": \"hwchase17\",\n      \"enabled\": true,\n      \"priority\": true,\n      \"topics\": [\n        \"ai-agent\"\n      ],\n      \"note\": \"LangChain founder, Agent ecosystem\"\n    },\n    {\n      \"id\": \"swyx-twitter\",\n      \"type\": \"twitter\",\n      \"name\": \"Swyx\",\n      \"handle\": \"swyx\",\n      \"enabled\": true,\n      \"priority\": false,\n      \"topics\": [\n        \"llm\",\n        \"ai-agent\"\n      ],\n      \"note\": \"AI Engineer community, Latent Space podcast\"\n    },\n    {\n      \"id\": \"erikbryn-twitter\",\n      \"type\": \"twitter\",\n      \"name\": \"Erik Brynjolfsson\",\n      \"handle\": \"erikbryn\",\n      \"enabled\": true,\n      \"priority\": false,\n      \"topics\": [\n        \"frontier-tech\"\n      ],\n      \"note\": \"Stanford Digital Economy Lab\"\n    },\n    {\n      \"id\": \"balaji-twitter\",\n      \"type\": \"twitter\",\n      \"name\": \"Balaji Srinivasan\",\n      \"handle\": \"balaji\",\n      \"enabled\": true,\n      \"priority\": true,\n      \"topics\": [\n        \"crypto\",\n        \"frontier-tech\"\n      ],\n      \"note\": \"Former Coinbase CTO, macro thinker\"\n    },\n    {\n      \"id\": \"cobie-twitter\",\n      \"type\": \"twitter\",\n      \"name\": \"Cobie\",\n      \"handle\": \"cobie\",\n      \"enabled\": true,\n      \"priority\": false,\n      \"topics\": [\n        \"crypto\"\n      ],\n      \"note\": \"Influential independent crypto analyst\"\n    },\n    {\n      \"id\": \"hsaka-twitter\",\n      \"type\": \"twitter\",\n      \"name\": \"Hsaka\",\n      \"handle\": \"HsakaTrades\",\n      \"enabled\": true,\n      \"priority\": false,\n      \"topics\": [\n        \"crypto\"\n      ],\n      \"note\": \"Crypto trading analysis\"\n    },\n    {\n      \"id\": \"cochran-twitter\",\n      \"type\": \"twitter\",\n      \"name\": \"Adam Cochran\",\n      \"handle\": \"adamscochran\",\n      \"enabled\": true,\n      \"priority\": false,\n      \"topics\": [\n        \"crypto\"\n      ],\n      \"note\": \"Cinneamhain Ventures, on-chain analysis\"\n    },\n    {\n      \"id\": \"lcermak-twitter\",\n      \"type\": \"twitter\",\n      \"name\": \"Larry Cermak\",\n      \"handle\": \"lawmaster\",\n      \"enabled\": true,\n      \"priority\": false,\n      \"topics\": [\n        \"crypto\"\n      ],\n      \"note\": \"The Block research director\"\n    },\n    {\n      \"id\": \"pytorch-github\",\n      \"type\": \"github\",\n      \"name\": \"PyTorch\",\n      \"repo\": \"pytorch/pytorch\",\n      \"enabled\": true,\n      \"priority\": true,\n      \"topics\": [\n        \"llm\",\n        \"frontier-tech\"\n      ],\n      \"note\": \"Leading deep learning framework from Meta\"\n    },\n    {\n      \"id\": \"transformers-github\",\n      \"type\": \"github\",\n      \"name\": \"Hugging Face Transformers\",\n      \"repo\": \"huggingface/transformers\",\n      \"enabled\": true,\n      \"priority\": true,\n      \"topics\": [\n        \"llm\"\n      ],\n      \"note\": \"State-of-the-art transformer models library\"\n    },\n    {\n      \"id\": \"langchain-github\",\n      \"type\": \"github\",\n      \"name\": \"LangChain\",\n      \"repo\": \"langchain-ai/langchain\",\n      \"enabled\": true,\n      \"priority\": true,\n      \"topics\": [\n        \"ai-agent\"\n      ],\n      \"note\": \"Framework for building LLM applications\"\n    },\n    {\n      \"id\": \"llamaindex-github\",\n      \"type\": \"github\",\n      \"name\": \"LlamaIndex\",\n      \"repo\": \"run-llama/llama_index\",\n      \"enabled\": true,\n      \"priority\": false,\n      \"topics\": [\n        \"ai-agent\"\n      ],\n      \"note\": \"Data framework for LLM applications\"\n    },\n    {\n      \"id\": \"ollama-github\",\n      \"type\": \"github\",\n      \"name\": \"Ollama\",\n      \"repo\": \"ollama/ollama\",\n      \"enabled\": true,\n      \"priority\": true,\n      \"topics\": [\n        \"llm\"\n      ],\n      \"note\": \"Run LLMs locally with ease\"\n    },\n    {\n      \"id\": \"vllm-github\",\n      \"type\": \"github\",\n      \"name\": \"vLLM\",\n      \"repo\": \"vllm-project/vllm\",\n      \"enabled\": true,\n      \"priority\": false,\n      \"topics\": [\n        \"llm\"\n      ],\n      \"note\": \"High-throughput LLM inference engine\"\n    },\n    {\n      \"id\": \"openai-python-github\",\n      \"type\": \"github\",\n      \"name\": \"OpenAI Python SDK\",\n      \"repo\": \"openai/openai-python\",\n      \"enabled\": true,\n      \"priority\": false,\n      \"topics\": [\n        \"llm\"\n      ],\n      \"note\": \"Official OpenAI Python client library\"\n    },\n    {\n      \"id\": \"anthropic-sdk-github\",\n      \"type\": \"github\",\n      \"name\": \"Anthropic SDK\",\n      \"repo\": \"anthropics/anthropic-sdk-python\",\n      \"enabled\": true,\n      \"priority\": false,\n      \"topics\": [\n        \"llm\"\n      ],\n      \"note\": \"Official Anthropic Python SDK\"\n    },\n    {\n      \"id\": \"crewai-github\",\n      \"type\": \"github\",\n      \"name\": \"CrewAI\",\n      \"repo\": \"crewAIInc/crewAI\",\n      \"enabled\": true,\n      \"priority\": false,\n      \"topics\": [\n        \"ai-agent\"\n      ],\n      \"note\": \"Multi-agent AI collaboration framework\"\n    },\n    {\n      \"id\": \"autogen-github\",\n      \"type\": \"github\",\n      \"name\": \"AutoGen\",\n      \"repo\": \"microsoft/autogen\",\n      \"enabled\": true,\n      \"priority\": false,\n      \"topics\": [\n        \"ai-agent\"\n      ],\n      \"note\": \"Multi-agent conversation framework from Microsoft\"\n    },\n    {\n      \"id\": \"dify-github\",\n      \"type\": \"github\",\n      \"name\": \"Dify\",\n      \"repo\": \"langgenius/dify\",\n      \"enabled\": true,\n      \"priority\": false,\n      \"topics\": [\n        \"ai-agent\"\n      ],\n      \"note\": \"LLM app development platform\"\n    },\n    {\n      \"id\": \"openclaw-github\",\n      \"type\": \"github\",\n      \"name\": \"OpenClaw\",\n      \"repo\": \"openclaw/openclaw\",\n      \"enabled\": true,\n      \"priority\": true,\n      \"topics\": [\n        \"ai-agent\"\n      ],\n      \"note\": \"Open source AI assistant platform\"\n    },\n    {\n      \"id\": \"go-ethereum-github\",\n      \"type\": \"github\",\n      \"name\": \"go-ethereum (Geth)\",\n      \"repo\": \"ethereum/go-ethereum\",\n      \"enabled\": true,\n      \"priority\": true,\n      \"topics\": [\n        \"crypto\"\n      ],\n      \"note\": \"Official Go implementation of Ethereum\"\n    },\n    {\n      \"id\": \"solidity-github\",\n      \"type\": \"github\",\n      \"name\": \"Solidity\",\n      \"repo\": \"ethereum/solidity\",\n      \"enabled\": true,\n      \"priority\": false,\n      \"topics\": [\n        \"crypto\"\n      ],\n      \"note\": \"Ethereum smart contract programming language\"\n    },\n    {\n      \"id\": \"foundry-github\",\n      \"type\": \"github\",\n      \"name\": \"Foundry\",\n      \"repo\": \"foundry-rs/foundry\",\n      \"enabled\": true,\n      \"priority\": false,\n      \"topics\": [\n        \"crypto\"\n      ],\n      \"note\": \"Fast, portable and modular Ethereum toolkit\"\n    },\n    {\n      \"id\": \"eips-github\",\n      \"type\": \"github\",\n      \"name\": \"Ethereum EIPs\",\n      \"repo\": \"ethereum/EIPs\",\n      \"enabled\": true,\n      \"priority\": true,\n      \"topics\": [\n        \"crypto\"\n      ],\n      \"note\": \"Ethereum Improvement Proposals repository\"\n    },\n    {\n      \"id\": \"linux-github\",\n      \"type\": \"github\",\n      \"name\": \"Linux Kernel\",\n      \"repo\": \"torvalds/linux\",\n      \"enabled\": true,\n      \"priority\": false,\n      \"topics\": [\n        \"frontier-tech\"\n      ],\n      \"note\": \"Linux kernel source code maintained by Linus Torvalds\"\n    },\n    {\n      \"id\": \"rust-github\",\n      \"type\": \"github\",\n      \"name\": \"Rust\",\n      \"repo\": \"rust-lang/rust\",\n      \"enabled\": true,\n      \"priority\": false,\n      \"topics\": [\n        \"frontier-tech\"\n      ],\n      \"note\": \"The Rust programming language compiler and standard library\"\n    },\n    {\n      \"id\": \"agno-github\",\n      \"type\": \"github\",\n      \"name\": \"Agno\",\n      \"repo\": \"agno-agi/agno\",\n      \"enabled\": true,\n      \"priority\": true,\n      \"topics\": [\n        \"ai-agent\"\n      ],\n      \"note\": \"Lightweight AI agent framework\"\n    },\n    {\n      \"id\": \"bensbites-rss\",\n      \"type\": \"rss\",\n      \"name\": \"Ben's Bites\",\n      \"url\": \"https://www.bensbites.com/feed\",\n      \"enabled\": true,\n      \"priority\": true,\n      \"topics\": [\n        \"llm\",\n        \"ai-agent\"\n      ]\n    },\n    {\n      \"id\": \"the-decoder-rss\",\n      \"type\": \"rss\",\n      \"name\": \"The Decoder\",\n      \"url\": \"https://the-decoder.com/feed/\",\n      \"enabled\": true,\n      \"priority\": false,\n      \"topics\": [\n        \"llm\",\n        \"ai-agent\"\n      ]\n    },\n    {\n      \"id\": \"a16zcrypto-rss\",\n      \"type\": \"rss\",\n      \"name\": \"a16z Crypto\",\n      \"url\": \"https://a16zcrypto.substack.com/feed\",\n      \"enabled\": true,\n      \"priority\": true,\n      \"topics\": [\n        \"crypto\"\n      ]\n    },\n    {\n      \"id\": \"bankless-rss\",\n      \"type\": \"rss\",\n      \"name\": \"Bankless\",\n      \"url\": \"https://newsletter.banklesshq.com/feed\",\n      \"enabled\": true,\n      \"priority\": true,\n      \"topics\": [\n        \"crypto\"\n      ]\n    },\n    {\n      \"id\": \"twitter-clementdelangue\",\n      \"type\": \"twitter\",\n      \"name\": \"Clement Delangue\",\n      \"handle\": \"ClementDelangue\",\n      \"enabled\": true,\n      \"priority\": true,\n      \"topics\": [\n        \"llm\",\n        \"ai-agent\"\n      ]\n    },\n    {\n      \"id\": \"twitter-gaborhm\",\n      \"type\": \"twitter\",\n      \"name\": \"Greg Brockman\",\n      \"handle\": \"GregBrockman\",\n      \"enabled\": true,\n      \"priority\": true,\n      \"topics\": [\n        \"llm\",\n        \"ai-agent\"\n      ]\n    },\n    {\n      \"id\": \"twitter-zuck\",\n      \"type\": \"twitter\",\n      \"name\": \"Mark Zuckerberg\",\n      \"handle\": \"finkd\",\n      \"enabled\": true,\n      \"priority\": true,\n      \"topics\": [\n        \"llm\",\n        \"ai-agent\"\n      ]\n    },\n    {\n      \"id\": \"github-mcp-servers\",\n      \"type\": \"github\",\n      \"name\": \"MCP Servers\",\n      \"repo\": \"modelcontextprotocol/servers\",\n      \"enabled\": true,\n      \"priority\": true,\n      \"topics\": [\n        \"ai-agent\"\n      ]\n    },\n    {\n      \"id\": \"github-deepseek-v3\",\n      \"type\": \"github\",\n      \"name\": \"DeepSeek V3\",\n      \"repo\": \"deepseek-ai/DeepSeek-V3\",\n      \"enabled\": true,\n      \"priority\": true,\n      \"topics\": [\n        \"llm\"\n      ]\n    },\n    {\n      \"id\": \"github-meta-llama\",\n      \"type\": \"github\",\n      \"name\": \"Meta Llama\",\n      \"repo\": \"meta-llama/llama-models\",\n      \"enabled\": true,\n      \"priority\": true,\n      \"topics\": [\n        \"llm\"\n      ]\n    },\n    {\n      \"id\": \"reddit-machinelearning\",\n      \"type\": \"reddit\",\n      \"name\": \"r/MachineLearning\",\n      \"subreddit\": \"MachineLearning\",\n      \"sort\": \"hot\",\n      \"limit\": 25,\n      \"min_score\": 50,\n      \"enabled\": true,\n      \"priority\": true,\n      \"topics\": [\n        \"llm\",\n        \"ai-agent\"\n      ]\n    },\n    {\n      \"id\": \"reddit-localllama\",\n      \"type\": \"reddit\",\n      \"name\": \"r/LocalLLaMA\",\n      \"subreddit\": \"LocalLLaMA\",\n      \"sort\": \"hot\",\n      \"limit\": 25,\n      \"min_score\": 30,\n      \"enabled\": true,\n      \"priority\": false,\n      \"topics\": [\n        \"llm\"\n      ]\n    },\n    {\n      \"id\": \"reddit-cryptocurrency\",\n      \"type\": \"reddit\",\n      \"name\": \"r/CryptoCurrency\",\n      \"subreddit\": \"CryptoCurrency\",\n      \"sort\": \"hot\",\n      \"limit\": 25,\n      \"min_score\": 50,\n      \"enabled\": true,\n      \"priority\": true,\n      \"topics\": [\n        \"crypto\"\n      ]\n    },\n    {\n      \"id\": \"reddit-artificial\",\n      \"type\": \"reddit\",\n      \"name\": \"r/artificial\",\n      \"subreddit\": \"artificial\",\n      \"sort\": \"hot\",\n      \"limit\": 25,\n      \"min_score\": 30,\n      \"enabled\": true,\n      \"priority\": false,\n      \"topics\": [\n        \"llm\",\n        \"ai-agent\"\n      ]\n    },\n    {\n      \"id\": \"reddit-ethereum\",\n      \"type\": \"reddit\",\n      \"name\": \"r/ethereum\",\n      \"subreddit\": \"ethereum\",\n      \"sort\": \"hot\",\n      \"limit\": 25,\n      \"min_score\": 30,\n      \"enabled\": true,\n      \"priority\": false,\n      \"topics\": [\n        \"crypto\"\n      ]\n    },\n    {\n      \"id\": \"reddit-chatgpt\",\n      \"type\": \"reddit\",\n      \"name\": \"r/ChatGPT\",\n      \"subreddit\": \"ChatGPT\",\n      \"sort\": \"hot\",\n      \"limit\": 25,\n      \"min_score\": 50,\n      \"enabled\": true,\n      \"priority\": true,\n      \"topics\": [\n        \"llm\",\n        \"ai-agent\"\n      ]\n    },\n    {\n      \"id\": \"reddit-singularity\",\n      \"type\": \"reddit\",\n      \"name\": \"r/singularity\",\n      \"subreddit\": \"singularity\",\n      \"sort\": \"hot\",\n      \"limit\": 25,\n      \"min_score\": 50,\n      \"enabled\": true,\n      \"priority\": false,\n      \"topics\": [\n        \"llm\"\n      ]\n    },\n    {\n      \"id\": \"reddit-openai\",\n      \"type\": \"reddit\",\n      \"name\": \"r/OpenAI\",\n      \"subreddit\": \"OpenAI\",\n      \"sort\": \"hot\",\n      \"limit\": 25,\n      \"min_score\": 50,\n      \"enabled\": true,\n      \"priority\": true,\n      \"topics\": [\n        \"llm\",\n        \"ai-agent\"\n      ]\n    },\n    {\n      \"id\": \"reddit-bitcoin\",\n      \"type\": \"reddit\",\n      \"name\": \"r/Bitcoin\",\n      \"subreddit\": \"Bitcoin\",\n      \"sort\": \"hot\",\n      \"limit\": 25,\n      \"min_score\": 50,\n      \"enabled\": true,\n      \"priority\": true,\n      \"topics\": [\n        \"crypto\"\n      ]\n    },\n    {\n      \"id\": \"reddit-programming\",\n      \"type\": \"reddit\",\n      \"name\": \"r/programming\",\n      \"subreddit\": \"programming\",\n      \"sort\": \"hot\",\n      \"limit\": 25,\n      \"min_score\": 50,\n      \"enabled\": true,\n      \"priority\": false,\n      \"topics\": [\n        \"llm\",\n        \"ai-agent\",\n        \"crypto\"\n      ]\n    },\n    {\n      \"id\": \"reddit-anthropic\",\n      \"type\": \"reddit\",\n      \"name\": \"r/Anthropic\",\n      \"subreddit\": \"Anthropic\",\n      \"sort\": \"hot\",\n      \"limit\": 25,\n      \"min_score\": 30,\n      \"enabled\": true,\n      \"priority\": true,\n      \"topics\": [\n        \"llm\",\n        \"ai-agent\"\n      ]\n    },\n    {\n      \"id\": \"reddit-defi\",\n      \"type\": \"reddit\",\n      \"name\": \"r/defi\",\n      \"subreddit\": \"defi\",\n      \"sort\": \"hot\",\n      \"limit\": 25,\n      \"min_score\": 30,\n      \"enabled\": true,\n      \"priority\": false,\n      \"topics\": [\n        \"crypto\"\n      ]\n    },\n    {\n      \"id\": \"reddit-experienceddevs\",\n      \"type\": \"reddit\",\n      \"name\": \"r/ExperiencedDevs\",\n      \"subreddit\": \"ExperiencedDevs\",\n      \"sort\": \"hot\",\n      \"limit\": 25,\n      \"min_score\": 50,\n      \"enabled\": true,\n      \"priority\": false,\n      \"topics\": [\n        \"llm\",\n        \"ai-agent\"\n      ]\n    },\n    {\n      \"id\": \"openclaw-twitter\",\n      \"type\": \"twitter\",\n      \"name\": \"OpenClaw\",\n      \"handle\": \"OpenClawAI\",\n      \"enabled\": true,\n      \"priority\": true,\n      \"topics\": [\n        \"ai-agent\"\n      ],\n      \"note\": \"Official OpenClaw account\"\n    },\n    {\n      \"id\": \"steipete-twitter\",\n      \"type\": \"twitter\",\n      \"name\": \"Peter Steinberger\",\n      \"handle\": \"steipete\",\n      \"enabled\": true,\n      \"priority\": true,\n      \"topics\": [\n        \"ai-agent\",\n        \"frontier-tech\"\n      ],\n      \"note\": \"OpenClaw creator, now at OpenAI\"\n    },\n    {\n      \"id\": \"mem0-github\",\n      \"type\": \"github\",\n      \"repo\": \"mem0ai/mem0\",\n      \"name\": \"Mem0\",\n      \"enabled\": true,\n      \"priority\": false,\n      \"topics\": [\n        \"ai-agent\"\n      ],\n      \"note\": \"Memory layer for AI agents\"\n    },\n    {\n      \"id\": \"openviking-github\",\n      \"type\": \"github\",\n      \"repo\": \"volcengine/OpenViking\",\n      \"name\": \"OpenViking\",\n      \"enabled\": true,\n      \"priority\": false,\n      \"topics\": [\n        \"llm\"\n      ],\n      \"note\": \"Volcengine/ByteDance open-source LLM\"\n    },\n    {\n      \"id\": \"moltworker-github\",\n      \"type\": \"github\",\n      \"name\": \"Cloudflare MoltWorker\",\n      \"repo\": \"cloudflare/moltworker\",\n      \"enabled\": true,\n      \"priority\": false,\n      \"topics\": [\n        \"frontier-tech\"\n      ],\n      \"note\": \"Cloudflare MoltWorker project\"\n    },\n    {\n      \"id\": \"picoclaw-github\",\n      \"type\": \"github\",\n      \"name\": \"Sipeed PicoClaw\",\n      \"repo\": \"sipeed/picoclaw\",\n      \"enabled\": true,\n      \"priority\": false,\n      \"topics\": [\n        \"frontier-tech\"\n      ],\n      \"note\": \"Sipeed PicoClaw embedded AI project\"\n    },\n    {\n      \"id\": \"nanobot-github\",\n      \"type\": \"github\",\n      \"name\": \"HKUDS NanoBot\",\n      \"repo\": \"HKUDS/nanobot\",\n      \"enabled\": true,\n      \"priority\": false,\n      \"topics\": [\n        \"ai-agent\"\n      ],\n      \"note\": \"HKUDS NanoBot AI agent project\"\n    },\n    {\n      \"id\": \"zeroclaw-github\",\n      \"type\": \"github\",\n      \"name\": \"ZeroClaw\",\n      \"repo\": \"zeroclaw-labs/zeroclaw\",\n      \"enabled\": true,\n      \"priority\": false,\n      \"topics\": [\n        \"crypto\"\n      ],\n      \"note\": \"ZeroClaw Labs project\"\n    }\n  ]\n}\n\nFile v3.11.0:config/defaults/topics.json\n\n{\n  \"_description\": \"Enhanced topic definitions for tech digest. Each topic defines a report section with search queries, filters, and display preferences.\",\n  \"_updated\": \"2026-02-15\",\n  \"_version\": \"2.5.0\",\n  \"topics\": [\n    {\n      \"id\": \"llm\",\n      \"emoji\": \"🧠\",\n      \"label\": \"LLM / Large Models\",\n      \"description\": \"Large Language Models, foundation models, model releases, benchmarks, and breakthroughs in generative AI\",\n      \"search\": {\n        \"queries\": [\"LLM latest news\", \"large language model breakthroughs\", \"大模型最新动态\", \"GPT Claude Gemini updates\"],\n        \"twitter_queries\": [\"GPT-5\", \"Claude\", \"大模型\"],\n        \"must_include\": [\"LLM\", \"large language model\", \"foundation model\", \"language model\", \"大模型\"],\n        \"exclude\": [\"tutorial\", \"how to use\", \"beginner guide\"]\n      },\n      \"display\": {\n        \"max_items\": 8,\n        \"style\": \"detailed\"\n      }\n    },\n    {\n      \"id\": \"ai-agent\",\n      \"emoji\": \"🤖\",\n      \"label\": \"AI Agent\",\n      \"description\": \"Autonomous agents, agent frameworks, AI assistants, and agentic AI systems\",\n      \"search\": {\n        \"queries\": [\"AI Agent latest developments\", \"autonomous agent framework\", \"AI assistant breakthrough\"],\n        \"twitter_queries\": [\"AI agent\", \"autonomous agent\", \"AI 智能体\"],\n        \"must_include\": [\"AI agent\", \"autonomous agent\", \"agent framework\", \"agentic\", \"multi-agent\"],\n        \"exclude\": [\"game agent\", \"travel agent\"]\n      },\n      \"display\": {\n        \"max_items\": 6,\n        \"style\": \"compact\"\n      }\n    },\n    {\n      \"id\": \"crypto\",\n      \"emoji\": \"💰\",\n      \"label\": \"Cryptocurrency\",\n      \"description\": \"Bitcoin, Ethereum, DeFi, NFTs, blockchain technology, and crypto market developments\",\n      \"search\": {\n        \"queries\": [\"cryptocurrency bitcoin ethereum latest news\", \"加密货币最新新闻\", \"DeFi breakthrough\", \"blockchain development\"],\n        \"twitter_queries\": [\"Bitcoin\", \"Ethereum\", \"加密货币\"],\n        \"must_include\": [\"crypto\", \"bitcoin\", \"ethereum\", \"blockchain\", \"DeFi\", \"NFT\", \"web3\"],\n        \"exclude\": [\"scam\", \"pump dump\", \"get rich quick\"]\n      },\n      \"display\": {\n        \"max_items\": 6,\n        \"style\": \"compact\"\n      }\n    },\n    {\n      \"id\": \"frontier-tech\",\n      \"emoji\": \"🔬\",\n      \"label\": \"Frontier Tech\",\n      \"description\": \"Cutting-edge technology, research breakthroughs, quantum computing, biotech, and emerging technologies\",\n      \"search\": {\n        \"queries\": [\"artificial intelligence breakthroughs\", \"frontier technology latest\", \"quantum computing progress\", \"biotech breakthrough\"],\n        \"twitter_queries\": [\"AI breakthrough\", \"量子计算\", \"机器人\"],\n        \"must_include\": [\"breakthrough\", \"research\", \"technology\", \"innovation\", \"quantum\", \"biotech\", \"robotics\"],\n        \"exclude\": [\"rumor\", \"speculation\", \"unverified\"]\n      },\n      \"display\": {\n        \"max_items\": 8,\n        \"style\": \"detailed\"\n      }\n    }\n  ]\n}\n\nArchive v3.10.3: 29 files, 96907 bytes\n\nFiles: CHANGELOG.md (16948b), config/defaults/sources.json (38686b), config/defaults/topics.json (2951b), config/schema.json (4561b), CONTRIBUTING.md (2283b), README_CN.md (2950b), README.md (3198b), references/digest-prompt.md (5549b), references/templates/discord.md (2879b), references/templates/email.md (5783b), references/templates/pdf.md (2119b), requirements.txt (606b), scripts/config_loader.py (8490b), scripts/fetch-github.py (20476b), scripts/fetch-reddit.py (12211b), scripts/fetch-rss.py (19352b), scripts/fetch-twitter.py (28449b), scripts/fetch-web.py (20565b), scripts/generate-pdf.py (9751b), scripts/merge-sources.py (24247b), scripts/run-pipeline.py (9960b), scripts/sanitize-html.py (7357b), scripts/send-email.py (5199b), scripts/source-health.py (4931b), scripts/summarize-merged.py (3526b), scripts/test-pipeline.sh (10975b), scripts/validate-config.py (9659b), SKILL.md (20410b), _meta.json (136b)\n\nFile v3.10.3:SKILL.md\n\n---\nname: tech-news-digest\ndescription: Generate tech news digests with unified source model, quality scoring, and multi-format output. Five-layer data collection from RSS feeds, Twitter/X KOLs, GitHub releases, Reddit, and web search. Pipeline-based scripts with retry mechanisms and deduplication. Supports Discord, email, and markdown templates.\nversion: \"3.10.3\"\nhomepage: https://github.com/draco-agent/tech-news-digest\nsource: https://github.com/draco-agent/tech-news-digest\nmetadata:\n  openclaw:\n    requires:\n      bins: [\"python3\"]\n    optionalBins: [\"mail\", \"msmtp\", \"gog\", \"gh\", \"openssl\", \"weasyprint\"]\nenv:\n  - name: TWITTER_API_BACKEND\n    required: false\n    description: \"Twitter API backend: 'official', 'twitterapiio', or 'auto' (default: auto)\"\n  - name: X_BEARER_TOKEN\n    required: false\n    description: Twitter/X API bearer token for KOL monitoring (official backend)\n  - name: TWITTERAPI_IO_KEY\n    required: false\n    description: twitterapi.io API key for KOL monitoring (twitterapiio backend)\n  - name: BRAVE_API_KEYS\n    required: false\n    description: Brave Search API keys (comma-separated for rotation)\n  - name: BRAVE_API_KEY\n    required: false\n    description: Brave Search API key (single key fallback)\n  - name: GITHUB_TOKEN\n    required: false\n    description: GitHub token for higher API rate limits (auto-generated from GitHub App if not set)\n  - name: GH_APP_ID\n    required: false\n    description: GitHub App ID for automatic installation token generation\n  - name: GH_APP_INSTALL_ID\n    required: false\n    description: GitHub App Installation ID for automatic token generation\n  - name: GH_APP_KEY_FILE\n    required: false\n    description: Path to GitHub App private key PEM file\ntools:\n  - python3: Required. Runs data collection and merge scripts.\n  - mail: Optional. msmtp-based mail command for email delivery (preferred).\n  - gog: Optional. Gmail CLI for email delivery (fallback if mail not available).\nfiles:\n  read:\n    - config/defaults/: Default source and topic configurations\n    - references/: Prompt templates and output templates\n    - scripts/: Python pipeline scripts\n    - <workspace>/archive/tech-news-digest/: Previous digests for dedup\n  write:\n    - /tmp/td-*.json: Temporary pipeline intermediate outputs\n    - /tmp/td-email.html: Temporary email HTML body\n    - /tmp/td-digest.pdf: Generated PDF digest\n    - <workspace>/archive/tech-news-digest/: Saved digest archives\n---\n\n# Tech News Digest\n\nAutomated tech news digest system with unified data source model, quality scoring pipeline, and template-based output generation.\n\n## Quick Start\n\n1. **Configuration Setup**: Default configs are in `config/defaults/`. Copy to workspace for customization:\n   ```bash\n   mkdir -p workspace/config\n   cp config/defaults/sources.json workspace/config/tech-news-digest-sources.json\n   cp config/defaults/topics.json workspace/config/tech-news-digest-topics.json\n   ```\n\n2. **Environment Variables**: \n   - `TWITTERAPI_IO_KEY` - twitterapi.io API key (optional, preferred)\n   - `X_BEARER_TOKEN` - Twitter/X official API bearer token (optional, fallback)\n   - `BRAVE_API_KEYS` - Brave Search API keys, comma-separated for rotation (optional)\n   - `BRAVE_API_KEY` - Single Brave key fallback (optional)\n   - `GITHUB_TOKEN` - GitHub personal access token (optional, improves rate limits)\n\n3. **Generate Digest**:\n   ```bash\n   # Unified pipeline (recommended) — runs all 5 sources in parallel + merge\n   python3 scripts/run-pipeline.py \\\n     --defaults config/defaults \\\n     --config workspace/config \\\n     --hours 48 --freshness pd \\\n     --archive-dir workspace/archive/tech-news-digest/ \\\n     --output /tmp/td-merged.json --verbose --force\n   ```\n\n4. **Use Templates**: Apply Discord, email, or PDF templates to merged output\n\n## Configuration Files\n\n### `sources.json` - Unified Data Sources\n```json\n{\n  \"sources\": [\n    {\n      \"id\": \"openai-rss\",\n      \"type\": \"rss\",\n      \"name\": \"OpenAI Blog\",\n      \"url\": \"https://openai.com/blog/rss.xml\",\n      \"enabled\": true,\n      \"priority\": true,\n      \"topics\": [\"llm\", \"ai-agent\"],\n      \"note\": \"Official OpenAI updates\"\n    },\n    {\n      \"id\": \"sama-twitter\",\n      \"type\": \"twitter\", \n      \"name\": \"Sam Altman\",\n      \"handle\": \"sama\",\n      \"enabled\": true,\n      \"priority\": true,\n      \"topics\": [\"llm\", \"frontier-tech\"],\n      \"note\": \"OpenAI CEO\"\n    }\n  ]\n}\n```\n\n### `topics.json` - Enhanced Topic Definitions\n```json\n{\n  \"topics\": [\n    {\n      \"id\": \"llm\",\n      \"emoji\": \"🧠\",\n      \"label\": \"LLM / Large Models\",\n      \"description\": \"Large Language Models, foundation models, breakthroughs\",\n      \"search\": {\n        \"queries\": [\"LLM latest news\", \"large language model breakthroughs\"],\n        \"must_include\": [\"LLM\", \"large language model\", \"foundation model\"],\n        \"exclude\": [\"tutorial\", \"beginner guide\"]\n      },\n      \"display\": {\n        \"max_items\": 8,\n        \"style\": \"detailed\"\n      }\n    }\n  ]\n}\n```\n\n## Scripts Pipeline\n\n### `run-pipeline.py` - Unified Pipeline (Recommended)\n```bash\npython3 scripts/run-pipeline.py \\\n  --defaults config/defaults [--config CONFIG_DIR] \\\n  --hours 48 --freshness pd \\\n  --archive-dir workspace/archive/tech-news-digest/ \\\n  --output /tmp/td-merged.json --verbose --force\n```\n- **Features**: Runs all 5 fetch steps in parallel, then merges + deduplicates + scores\n- **Output**: Final merged JSON ready for report generation (~30s total)\n- **Metadata**: Saves per-step timing and counts to `*.meta.json`\n- **GitHub Auth**: Auto-generates GitHub App token if `$GITHUB_TOKEN` not set\n- **Fallback**: If this fails, run individual scripts below\n\n### Individual Scripts (Fallback)\n\n#### `fetch-rss.py` - RSS Feed Fetcher\n```bash\npython3 scripts/fetch-rss.py [--defaults DIR] [--config DIR] [--hours 48] [--output FILE] [--verbose]\n```\n- Parallel fetching (10 workers), retry with backoff, feedparser + regex fallback\n- Timeout: 30s per feed, ETag/Last-Modified caching\n\n#### `fetch-twitter.py` - Twitter/X KOL Monitor\n```bash\npython3 scripts/fetch-twitter.py [--defaults DIR] [--config DIR] [--hours 48] [--output FILE] [--backend auto|official|twitterapiio]\n```\n- Backend auto-detection: uses twitterapi.io if `TWITTERAPI_IO_KEY` set, else official X API v2 if `X_BEARER_TOKEN` set\n- Rate limit handling, engagement metrics, retry with backoff\n\n#### `fetch-web.py` - Web Search Engine\n```bash\npython3 scripts/fetch-web.py [--defaults DIR] [--config DIR] [--freshness pd] [--output FILE]\n```\n- Auto-detects Brave API rate limit: paid plans → parallel queries, free → sequential\n- Without API: generates search interface for agents\n\n#### `fetch-github.py` - GitHub Releases Monitor\n```bash\npython3 scripts/fetch-github.py [--defaults DIR] [--config DIR] [--hours 168] [--output FILE]\n```\n- Parallel fetching (10 workers), 30s timeout\n- Auth priority: `$GITHUB_TOKEN` → GitHub App auto-generate → `gh` CLI → unauthenticated (60 req/hr)\n\n#### `fetch-reddit.py` - Reddit Posts Fetcher\n```bash\npython3 scripts/fetch-reddit.py [--defaults DIR] [--config DIR] [--hours 48] [--output FILE]\n```\n- Parallel fetching (4 workers), public JSON API (no auth required)\n- 13 subreddits with score filtering\n\n#### `merge-sources.py` - Quality Scoring & Deduplication\n```bash\npython3 scripts/merge-sources.py --rss FILE --twitter FILE --web FILE --github FILE --reddit FILE\n```\n- Quality scoring, title similarity dedup (85%), previous digest penalty\n- Output: topic-grouped articles sorted by score\n\n#### `validate-config.py` - Configuration Validator\n```bash\npython3 scripts/validate-config.py [--defaults DIR] [--config DIR] [--verbose]\n```\n- JSON schema validation, topic reference checks, duplicate ID detection\n\n#### `generate-pdf.py` - PDF Report Generator\n```bash\npython3 scripts/generate-pdf.py --input report.md --output digest.pdf [--verbose]\n```\n- Converts markdown digest to styled A4 PDF with Chinese typography (Noto Sans CJK SC)\n- Emoji icons, page headers/footers, blue accent theme. Requires `weasyprint`.\n\n#### `sanitize-html.py` - Safe HTML Email Converter\n```bash\npython3 scripts/sanitize-html.py --input report.md --output email.html [--verbose]\n```\n- Converts markdown to XSS-safe HTML email with inline CSS\n- URL whitelist (http/https only), HTML-escaped text content\n\n#### `source-health.py` - Source Health Monitor\n```bash\npython3 scripts/source-health.py --rss FILE --twitter FILE --github FILE --reddit FILE --web FILE [--verbose]\n```\n- Tracks per-source success/failure history over 7 days\n- Reports unhealthy sources (>50% failure rate)\n\n#### `summarize-merged.py` - Merged Data Summary\n```bash\npython3 scripts/summarize-merged.py --input merged.json [--top N] [--topic TOPIC]\n```\n- Human-readable summary of merged data for LLM consumption\n- Shows top articles per topic with scores and metrics\n\n## User Customization\n\n### Workspace Configuration Override\nPlace custom configs in `workspace/config/` to override defaults:\n\n- **Sources**: Append new sources, disable defaults with `\"enabled\": false`\n- **Topics**: Override topic definitions, search queries, display settings\n- **Merge Logic**: \n  - Sources with same `id` → user version takes precedence\n  - Sources with new `id` → appended to defaults\n  - Topics with same `id` → user version completely replaces default\n\n### Example Workspace Override\n```json\n// workspace/config/tech-news-digest-sources.json\n{\n  \"sources\": [\n    {\n      \"id\": \"simonwillison-rss\",\n      \"enabled\": false,\n      \"note\": \"Disabled: too noisy for my use case\"\n    },\n    {\n      \"id\": \"my-custom-blog\", \n      \"type\": \"rss\",\n      \"name\": \"My Custom Tech Blog\",\n      \"url\": \"https://myblog.com/rss\",\n      \"enabled\": true,\n      \"priority\": true,\n      \"topics\": [\"frontier-tech\"]\n    }\n  ]\n}\n```\n\n## Templates & Output\n\n### Discord Template (`references/templates/discord.md`)\n- Bullet list format with link suppression (`<link>`)\n- Mobile-optimized, emoji headers\n- 2000 character limit awareness\n\n### Email Template (`references/templates/email.md`) \n- Rich metadata, technical stats, archive links\n- Executive summary, top articles section\n- HTML-compatible formatting\n\n### PDF Template (`references/templates/pdf.md`)\n- A4 layout with Noto Sans CJK SC font for Chinese support\n- Emoji icons, page headers/footers with page numbers\n- Generated via `scripts/generate-pdf.py` (requires `weasyprint`)\n\n## Default Sources (138 total)\n\n- **RSS Feeds (49)**: AI labs, tech blogs, crypto news, Chinese tech media\n- **Twitter/X KOLs (48)**: AI researchers, crypto leaders, tech executives\n- **GitHub Repos (28)**: Major open-source projects (LangChain, vLLM, DeepSeek, Llama, etc.)\n- **Reddit (13)**: r/MachineLearning, r/LocalLLaMA, r/CryptoCurrency, r/ChatGPT, r/OpenAI, etc.\n- **Web Search (4 topics)**: LLM, AI Agent, Crypto, Frontier Tech\n\nAll sources pre-configured with appropriate topic tags and priority levels.\n\n## Dependencies\n\n```bash\npip install -r requirements.txt\n```\n\n**Optional but Recommended**:\n- `feedparser>=6.0.0` - Better RSS parsing (fallback to regex if unavailable)\n- `jsonschema>=4.0.0` - Configuration validation\n\n**All scripts work with Python 3.8+ standard library only.**\n\n## Monitoring & Operations\n\n### Health Checks\n```bash\n# Validate configuration\npython3 scripts/validate-config.py --verbose\n\n# Test RSS feeds\npython3 scripts/fetch-rss.py --hours 1 --verbose\n\n# Check Twitter API\npython3 scripts/fetch-twitter.py --hours 1 --verbose\n```\n\n### Archive Management\n- Digests automatically archived to `<workspace>/archive/tech-news-digest/`\n- Previous digest titles used for duplicate detection\n- Old archives cleaned automatically (90+ days)\n\n### Error Handling\n- **Network Failures**: Retry with exponential backoff\n- **Rate Limits**: Automatic retry with appropriate delays\n- **Invalid Content**: Graceful degradation, detailed logging\n- **Configuration Errors**: Schema validation with helpful messages\n\n## API Keys & Environment\n\nSet in `~/.zshenv` or similar:\n```bash\n# Twitter (at least one required for Twitter source)\nexport TWITTERAPI_IO_KEY=\"your_key\"        # twitterapi.io key (preferred)\nexport X_BEARER_TOKEN=\"your_bearer_token\"  # Official X API v2 (fallback)\nexport TWITTER_API_BACKEND=\"auto\"          # auto|twitterapiio|official (default: auto)\n\n# Brave Search (optional, enables web search layer)\nexport BRAVE_API_KEYS=\"key1,key2,key3\"     # Multiple keys, comma-separated rotation\nexport BRAVE_API_KEY=\"key1\"                # Single key fallback\nexport BRAVE_PLAN=\"free\"                   # Override rate limit detection: free|pro\n\n# GitHub (optional, improves rate limits)\nexport GITHUB_TOKEN=\"ghp_xxx\"              # PAT (simplest)\nexport GH_APP_ID=\"12345\"                   # Or use GitHub App for auto-token\nexport GH_APP_INSTALL_ID=\"67890\"\nexport GH_APP_KEY_FILE=\"/path/to/key.pem\"\n```\n\n- **Twitter**: `TWITTERAPI_IO_KEY` preferred ($3-5/mo); `X_BEARER_TOKEN` as fallback; `auto` mode tries twitterapiio first\n- **Brave Search**: Optional, fallback to agent web_search if unavailable\n- **GitHub**: Auto-generates token from GitHub App if PAT not set; unauthenticated fallback (60 req/hr)\n- **Reddit**: No API key needed (uses public JSON API)\n\n## Cron / Scheduled Task Integration\n\n### OpenClaw Cron (Recommended)\n\nThe cron prompt should **NOT** hardcode the pipeline steps. Instead, reference `references/digest-prompt.md` and only pass configuration parameters. This ensures the pipeline logic stays in the skill repo and is consistent across all installations.\n\n#### Daily Digest Cron Prompt\n```\nRead <SKILL_DIR>/references/digest-prompt.md and follow the complete workflow to generate a daily digest.\n\nReplace placeholders with:\n- MODE = daily\n- TIME_WINDOW = past 1-2 days\n- FRESHNESS = pd\n- RSS_HOURS = 48\n- ITEMS_PER_SECTION = 3-5\n- BLOG_PICKS_COUNT = 2-3\n- EXTRA_SECTIONS = (none)\n- SUBJECT = Daily Tech Digest - YYYY-MM-DD\n- WORKSPACE = <your workspace path>\n- SKILL_DIR = <your skill install path>\n- DISCORD_CHANNEL_ID = <your channel id>\n- EMAIL = (optional)\n- LANGUAGE = English\n- TEMPLATE = discord\n\nFollow every step in the prompt template strictly. Do not skip any steps.\n```\n\n#### Weekly Digest Cron Prompt\n```\nRead <SKILL_DIR>/references/digest-prompt.md and follow the complete workflow to generate a weekly digest.\n\nReplace placeholders with:\n- MODE = weekly\n- TIME_WINDOW = past 7 days\n- FRESHNESS = pw\n- RSS_HOURS = 168\n- ITEMS_PER_SECTION = 5-8\n- BLOG_PICKS_COUNT = 3-5\n- EXTRA_SECTIONS = 📊 Weekly Trend Summary (2-3 sentences summarizing macro trends)\n- SUBJECT = Weekly Tech Digest - YYYY-MM-DD\n- WORKSPACE = <your workspace path>\n- SKILL_DIR = <your skill install path>\n- DISCORD_CHANNEL_ID = <your channel id>\n- EMAIL = (optional)\n- LANGUAGE = English\n- TEMPLATE = discord\n\nFollow every step in the prompt template strictly. Do not skip any steps.\n```\n\n#### Why This Pattern?\n- **Single source of truth**: Pipeline logic lives in `digest-prompt.md`, not scattered across cron configs\n- **Portable**: Same skill on different OpenClaw instances, just change paths and channel IDs\n- **Maintainable**: Update the skill → all cron jobs pick up changes automatically\n- **Anti-pattern**: Do NOT copy pipeline steps into the cron prompt — it will drift out of sync\n\n#### Multi-Channel Delivery Limitation\nOpenClaw enforces **cross-provider isolation**: a single session can only send messages to one provider (e.g., Discord OR Telegram, not both). If you need to deliver digests to multiple platforms, create **separate cron jobs** for each provider:\n\n```\n# Job 1: Discord + Email\n- DISCORD_CHANNEL_ID = <your-discord-channel-id>\n- EMAIL = user@example.com\n- TEMPLATE = discord\n\n# Job 2: Telegram DM\n- DISCORD_CHANNEL_ID = (none)\n- EMAIL = (none)\n- TEMPLATE = telegram\n```\nReplace `DISCORD_CHANNEL_ID` delivery with the target platform's delivery in the second job's prompt.\n\nThis is a security feature, not a bug — it prevents accidental cross-context data leakage.\n\n## Security Notes\n\n### Execution Model\nThis skill uses a **prompt template pattern**: the agent reads `digest-prompt.md` and follows its instructions. This is the standard OpenClaw skill execution model — the agent interprets structured instructions from skill-provided files. All instructions are shipped with the skill bundle and can be audited before installation.\n\n### Network Access\nThe Python scripts make outbound requests to:\n- RSS feed URLs (configured in `tech-news-digest-sources.json`)\n- Twitter/X API (`api.x.com` or `api.twitterapi.io`)\n- Brave Search API (`api.search.brave.com`)\n- GitHub API (`api.github.com`)\n- Reddit JSON API (`reddit.com`)\n\nNo data is sent to any other endpoints. All API keys are read from environment variables declared in the skill metadata.\n\n### Shell Safety\nEmail delivery uses `send-email.py` which constructs proper MIME multipart messages with HTML body + optional PDF attachment. Subject formats are hardcoded (`Daily Tech Digest - YYYY-MM-DD`). PDF generation uses `generate-pdf.py` via `weasyprint`. The prompt template explicitly prohibits interpolating untrusted content (article titles, tweet text, etc.) into shell arguments. Email addresses and subjects must be static placeholder values only.\n\n### File Access\nScripts read from `config/` and write to `workspace/archive/`. No files outside the workspace are accessed.\n\n## Support & Troubleshooting\n\n### Common Issues\n1. **RSS feeds failing**: Check network connectivity, use `--verbose` for details\n2. **Twitter rate limits**: Reduce sources or increase interval\n3. **Configuration errors**: Run `validate-config.py` for specific issues\n4. **No articles found**: Check time window (`--hours`) and source enablement\n\n### Debug Mode\nAll scripts support `--verbose` flag for detailed logging and troubleshooting.\n\n### Performance Tuning\n- **Parallel Workers**: Adjust `MAX_WORKERS` in scripts for your system\n- **Timeout Settings**: Increase `TIMEOUT` for slow networks\n- **Article Limits**: Adjust `MAX_ARTICLES_PER_FEED` based on needs\n## Security Considerations\n\n### Shell Execution\nThe digest prompt instructs agents to run Python scripts via shell commands. All script paths and arguments are skill-defined constants — no user input is interpolated into commands. Two scripts use `subprocess`:\n- `run-pipeline.py` orchestrates child fetch scripts (all within `scripts/` directory)\n- `fetch-github.py` has two subprocess calls:\n  1. `openssl dgst -sha256 -sign` for JWT signing (only if `GH_APP_*` env vars are set — signs a self-constructed JWT payload, no user content involved)\n  2. `gh auth token` CLI fallback (only if `gh` is installed — reads from gh's own credential store)\n\nNo user-supplied or fetched content is ever interpolated into subprocess arguments. Email delivery uses `send-email.py` which builds MIME messages programmatically — no shell interpolation. PDF generation uses `generate-pdf.py` via `weasyprint`. Email subjects are static format strings only — never constructed from fetched data.\n\n### Credential & File Access\nScripts do **not** directly read `~/.config/`, `~/.ssh/`, or any credential files. All API tokens are read from environment variables declared in the skill metadata. The GitHub auth cascade is:\n1. `$GITHUB_TOKEN` env var (you control what to provide)\n2. GitHub App token generation (only if you set `GH_APP_ID`, `GH_APP_INSTALL_ID`, and `GH_APP_KEY_FILE` — uses inline JWT signing via `openssl` CLI, no external scripts involved)\n3. `gh auth token` CLI (delegates to gh's own secure credential store)\n4. Unauthenticated (60 req/hr, safe fallback)\n\nIf you prefer no automatic credential discovery, simply set `$GITHUB_TOKEN` and the script will use it directly without attempting steps 2-3.\n\n### Dependency Installation\nThis skill does **not** install any packages. `requirements.txt` lists optional dependencies (`feedparser`, `jsonschema`) for reference only. All scripts work with Python 3.8+ standard library. Users should install optional deps in a virtualenv if desired — the skill never runs `pip install`.\n\n### Input Sanitization\n- URL resolution rejects non-HTTP(S) schemes (javascript:, data:, etc.)\n- RSS fallback parsing uses simple, non-backtracking regex patterns (no ReDoS risk)\n- All fetched content is treated as untrusted data for display only\n\n### Network Access\nScripts make outbound HTTP requests to configured RSS feeds, Twitter API, GitHub API, Reddit JSON API, and Brave Search API. No inbound connections or listeners are created.\n\nFile v3.10.3:README.md\n\n# Tech News Digest\n\n> Automated tech news digest — 138 sources, 5-layer pipeline, one chat message to install.\n\n[![Python 3.8+](https://img.shields.io/badge/python-3.8+-blue.svg)](https://www.python.org/downloads/)\n[![MIT License](https://img.shields.io/badge/license-MIT-green.svg)](LICENSE)\n\n## 💬 Install in One Message\n\nTell your [OpenClaw](https://openclaw.ai) AI assistant:\n\n> **\"Install tech-news-digest and send a daily digest to #tech-news every morning at 9am\"**\n\nThat's it. Your bot handles installation, configuration, scheduling, and delivery — all through conversation.\n\nMore examples:\n\n> 🗣️ \"Set up a weekly AI digest, only LLM and AI Agent topics, deliver to Discord #ai-weekly every Monday\"\n\n> 🗣️ \"Install tech-news-digest, add my RSS feeds, and send crypto news to Telegram\"\n\n> 🗣️ \"Give me a tech digest right now, skip Twitter sources\"\n\nOr install via CLI:\n```bash\nclawhub install tech-news-digest\n```\n\n## 📊 What You Get\n\nA quality-scored, deduplicated tech digest built from **138 sources**:\n\n| Layer | Sources | What |\n|-------|---------|------|\n| 📡 RSS | 49 feeds | OpenAI, Anthropic, Ben's Bites, HN, 36氪, CoinDesk… |\n| 🐦 Twitter/X | 48 KOLs | @karpathy, @VitalikButerin, @sama, @elonmusk… |\n| 🔍 Web Search | 4 topics | Brave Search API with freshness filters |\n| 🐙 GitHub | 28 repos | Releases from key projects (LangChain, vLLM, DeepSeek, Llama…) |\n| 🗣️ Reddit | 13 subs | r/MachineLearning, r/LocalLLaMA, r/CryptoCurrency… |\n\n### Pipeline\n\n```\n       run-pipeline.py (~30s)\n              ↓\n  RSS ─┐\n  Twitter ─┤\n  Web ─────┤── parallel fetch ──→ merge-sources.py\n  GitHub ──┤\n  Reddit ──┘\n              ↓\n  Quality Scoring → Deduplication → Topic Grouping\n              ↓\n    Discord / Email / PDF output\n```\n\n**Quality scoring**: priority source (+3), multi-source cross-ref (+5), recency (+2), engagement (+1), Reddit score bonus (+1/+3/+5), already reported (-5).\n\n## ⚙️ Configuration\n\n- `config/defaults/sources.json` — 138 built-in sources\n- `config/defaults/topics.json` — 4 topics with search queries & Twitter queries\n- User overrides in `workspace/config/` take priority\n\n## 🔧 Optional Setup\n\nAll environment variables are optional. The pipeline runs with whatever sources are available.\n\n```bash\nexport TWITTERAPI_IO_KEY=\"...\"  # twitterapi.io (~$5/mo) — enables Twitter layer\nexport X_BEARER_TOKEN=\"...\"     # Twitter/X official API — alternative Twitter backend\nexport BRAVE_API_KEYS=\"k1,k2,k3\" # Brave Search API keys (comma-separated, rotation)\nexport BRAVE_API_KEY=\"...\"       # Fallback: single Brave key\nexport GITHUB_TOKEN=\"...\"       # GitHub API — higher rate limits (auto-generated from GitHub App if unset)\nexport TWITTER_API_BACKEND=\"auto\" # auto|twitterapiio|official (default: auto)\nexport BRAVE_PLAN=\"free\"         # Override Brave rate limit detection: free|pro\npip install weasyprint           # Enables PDF report generation\n```\n\n## 📂 Repository\n\n**GitHub**: [github.com/draco-agent/tech-news-digest](https://github.com/draco-agent/tech-news-digest)\n\n## 📄 License\n\nMIT License — see [LICENSE](LICENSE) for details.\n\nFile v3.10.3:_meta.json\n\n{\n  \"ownerId\": \"kn74589cx1nbhnc3x0f3nwre39814699\",\n  \"slug\": \"tech-news-digest\",\n  \"version\": \"3.10.3\",\n  \"publishedAt\": 1772207168147\n}\n\nFile v3.10.3:references/digest-prompt.md\n\n# Digest Prompt Template\n\nReplace `<...>` placeholders before use. Daily defaults shown; weekly overrides in parentheses.\n\n## Placeholders\n\n| Placeholder | Default | Weekly Override |\n|-------------|---------|----------------|\n| `<MODE>` | `daily` | `weekly` |\n| `<TIME_WINDOW>` | `past 1-2 days` | `past 7 days` |\n| `<FRESHNESS>` | `pd` | `pw` |\n| `<RSS_HOURS>` | `48` | `168` |\n| `<ITEMS_PER_SECTION>` | `3-5` | `5-8` |\n| `<BLOG_PICKS_COUNT>` | `2-3` | `3-5` |\n| `<EXTRA_SECTIONS>` | *(none)* | `📊 Weekly Trend Summary` |\n| `<SUBJECT>` | `Daily Tech Digest - YYYY-MM-DD` | `Weekly Tech Digest - YYYY-MM-DD` |\n| `<WORKSPACE>` | Your workspace path | |\n| `<SKILL_DIR>` | Installed skill directory | |\n| `<DISCORD_CHANNEL_ID>` | Target channel ID | |\n| `<EMAIL>` | *(optional)* Recipient email | |\n| `<EMAIL_FROM>` | *(optional)* e.g. `MyBot <bot@example.com>` | |\n| `<LANGUAGE>` | `Chinese` | |\n| `<TEMPLATE>` | `discord` / `email` / `markdown` | |\n| `<DATE>` | Today's date YYYY-MM-DD (caller provides) | |\n| `<VERSION>` | Read from SKILL.md frontmatter | |\n\n---\n\nGenerate the <MODE> tech digest for **<DATE>**. Use `<DATE>` as the report date — do NOT infer it.\n\n## Configuration\n\nRead config files (workspace overrides take priority over defaults):\n1. **Sources**: `<WORKSPACE>/config/tech-news-digest-sources.json` → fallback `<SKILL_DIR>/config/defaults/sources.json`\n2. **Topics**: `<WORKSPACE>/config/tech-news-digest-topics.json` → fallback `<SKILL_DIR>/config/defaults/topics.json`\n\n## Context: Previous Report\n\nRead the most recent file from `<WORKSPACE>/archive/tech-news-digest/` to avoid repeats and follow up on developing stories. Skip if none exists.\n\n## Data Collection Pipeline\n\n**Use the unified pipeline** (runs all 5 sources in parallel, ~30s):\n\n```bash\npython3 <SKILL_DIR>/scripts/run-pipeline.py \\\n  --defaults <SKILL_DIR>/config/defaults \\\n  --config <WORKSPACE>/config \\\n  --hours <RSS_HOURS> --freshness <FRESHNESS> \\\n  --archive-dir <WORKSPACE>/archive/tech-news-digest/ \\\n  --output /tmp/td-merged.json --verbose --force\n```\n\nIf it fails, run individual scripts in `<SKILL_DIR>/scripts/` (see each script's `--help`), then merge with `merge-sources.py`.\n\n## Report Generation\n\nGet a structured overview:\n```bash\npython3 <SKILL_DIR>/scripts/summarize-merged.py --input /tmp/td-merged.json --top <ITEMS_PER_SECTION>\n```\n\nUse this output to select articles — **do NOT write ad-hoc Python to parse the JSON**. Apply the template from `<SKILL_DIR>/references/templates/<TEMPLATE>.md`.\n\nSelect articles **purely by quality_score regardless of source type**. For Reddit posts, append `*[Reddit r/xxx, {{score}}↑]*`.\n\n### Executive Summary\n2-4 sentences between title and topics, highlighting top 3-5 stories by score. Concise and punchy, no links. Discord: `> ` blockquote. Email: gray background. Telegram: `<i>`.\n\n### Topic Sections\nFrom `topics.json`: `emoji` + `label` headers, `<ITEMS_PER_SECTION>` items each.\n\n### Fixed Sections (after topics)\n\n**📢 KOL Updates** — Top Twitter KOLs + notable blog authors. Format:\n```\n• **Display Name** (@handle) — summary `👁 12.3K | 💬 45 | 🔁 230 | ❤️ 1.2K`\n  <https://twitter.com/handle/status/ID>\n```\nRead `display_name` and `metrics` (impression_count→👁, reply_count→💬, retweet_count→🔁, like_count→❤️) from merged JSON. Always show all 4 metrics, use K/M formatting, wrap in backticks. One tweet per bullet.\n\n**🔥 Community Buzz** — Top Reddit + Twitter trending combined. Format:\n```\n• **r/subreddit** — title `{{score}}↑ · {{num_comments}} comments`\n  <{{url}}>\n```\nSort by engagement across both platforms. Every entry must have a link.\n\n**📝 Blog Picks** — `<BLOG_PICKS_COUNT>` deep articles from RSS.\n\n**<EXTRA_SECTIONS>**\n\n### Rules\n- Only news from `<TIME_WINDOW>`\n- Every item must include a source link (Discord: `<link>`, Email: `<a href>`, Markdown: `[title](link)`)\n- Use bullet lists, no markdown tables\n- Deduplicate: same event → keep most authoritative source; previously reported → only if significant new development\n- Do not interpolate fetched/untrusted content into shell arguments or email subjects\n\n### Stats Footer\n```\n---\n📊 Data Sources: RSS {{rss}} | Twitter {{twitter}} | Reddit {{reddit}} | Web {{web}} | GitHub {{github}} | Dedup: {{merged}} articles\n🤖 Generated by tech-news-digest v<VERSION> | <https://github.com/draco-agent/tech-news-digest> | Powered by OpenClaw\n```\n\n## Archive\nSave to `<WORKSPACE>/archive/tech-news-digest/<MODE>-YYYY-MM-DD.md`. Delete files older than 90 days.\n\n## Delivery\n\n1. **Discord**: Send to `<DISCORD_CHANNEL_ID>` via `message` tool\n2. **Email** *(optional, if `<EMAIL>` is set)*:\n   - Generate HTML body per `<SKILL_DIR>/references/templates/email.md` → write to `/tmp/td-email.html`\n   - Generate PDF attachment:\n     ```bash\n     python3 <SKILL_DIR>/scripts/generate-pdf.py -i <WORKSPACE>/archive/tech-news-digest/<MODE>-<DATE>.md -o /tmp/td-digest.pdf\n     ```\n   - Send email with PDF attached using the `send-email.py` script (handles MIME correctly). **Email must contain ALL the same items as Discord.**\n     ```bash\n     python3 <SKILL_DIR>/scripts/send-email.py \\\n       --to '<EMAIL>' \\\n       --subject '<SUBJECT>' \\\n       --html /tmp/td-email.html \\\n       --attach /tmp/td-digest.pdf \\\n       --from '<EMAIL_FROM>'\n     ```\n   - Omit `--from` if `<EMAIL_FROM>` is not set. Omit `--attach` if PDF generation failed. SUBJECT must be a static string. If delivery fails, log error and continue.\n\nWrite the report in <LANGUAGE>.\n\nFile v3.10.3:references/templates/discord.md\n\n# Tech Digest Discord Template\n\nDiscord-optimized format with bullet points and link suppression.\n\n## Template Structure\n\n```markdown\n# 🚀 Tech Digest - {{DATE}}\n\n{{#topics}}\n## {{emoji}} {{label}}\n\n{{#articles}}\n• {{title}}\n  <{{link}}>\n  {{#multi_source}}*[{{source_count}} sources]*{{/multi_source}}\n\n{{/articles}}\n{{/topics}}\n\n---\n📊 Data Sources: RSS {{rss_count}} | Twitter {{twitter_count}} | Reddit {{reddit_count}} | Web {{web_count}} | GitHub {{github_count}} releases | After dedup: {{merged_count}} articles\n🤖 Generated by tech-news-digest v{{version}} | <https://github.com/draco-agent/tech-news-digest> | Powered by OpenClaw\n```\n\n## Delivery\n\n- **Default: Channel** — Send to the Discord channel specified by `DISCORD_CHANNEL_ID`\n- Use `message` tool with `target` set to the channel ID for channel delivery\n- For DM delivery instead, set `target` to a user ID\n\n## Discord-Specific Features\n\n- **Link suppression**: Wrap links in `<>` to prevent embeds\n- **Bullet format**: Use `•` for clean mobile display  \n- **No tables**: Discord mobile doesn't handle markdown tables well\n- **Emoji headers**: Visual hierarchy with topic emojis\n- **Concise metadata**: Source count and multi-source indicators\n- **Character limits**: Discord messages have 2000 char limit, may need splitting\n\n## Example Output\n\n```markdown\n# 🚀 Tech Digest - 2026-02-15\n\n## 🧠 LLM / Large Models\n\n• OpenAI releases GPT-5 with breakthrough reasoning capabilities\n  <https://openai.com/blog/gpt5-announcement>\n  *[3 sources]*\n\n• Meta's Llama 3.1 achieves new MMLU benchmarks\n  <https://ai.meta.com/blog/llama-31-release>\n\n## 🤖 AI Agent\n\n• LangChain launches production-ready agent framework\n  <https://blog.langchain.dev/production-agents>\n\n## 💰 Cryptocurrency\n\n• Bitcoin reaches new ATH at $67,000 amid ETF approval\n  <https://coindesk.com/markets/btc-ath-etf>\n  *[2 sources]*\n\n## 📢 KOL Updates\n\n• **Elon Musk** (@elonmusk) — Confirmed X's crypto trading feature `👁 2.1M | 💬 12.3K | 🔁 8.5K | ❤️ 49.8K`\n  <https://twitter.com/elonmusk/status/123456789>\n• **@saylor** — Valentine's BTC enthusiasm `👁 450K | 💬 1.2K | 🔁 3.1K | ❤️ 13K`\n  <https://twitter.com/saylor/status/987654321>\n\n---\n📊 Data Sources: RSS 285 | Twitter 67 | Reddit 45 | Web 60 | GitHub 29 releases | After dedup: 95 articles\n```\n\n## Variables\n\n- `{{DATE}}` - Report date (YYYY-MM-DD format)\n- `{{topics}}` - Array of topic objects\n- `{{emoji}}` - Topic emoji \n- `{{label}}` - Topic display name\n- `{{articles}}` - Array of article objects per topic\n- `{{title}}` - Article title (truncated if needed)\n- `{{link}}` - Article URL\n- `{{multi_source}}` - Boolean, true if article from multiple sources\n- `{{source_count}}` - Number of sources for this article\n- `{{total_sources}}` - Total number of sources used\n- `{{total_articles}}` - Total articles in digest\n\nFile v3.10.3:references/templates/email.md\n\n# Tech Digest Email Template\n\nHTML email format optimized for Gmail/Outlook rendering.\n\n## Delivery\n\nSend via `gog gmail send` with `--body-html` flag:\n```bash\ngog gmail send --to '<EMAIL>' --subject '<SUBJECT>' --body-html '<HTML_CONTENT>'\n```\n\n**Important**: Use `--body-html`, NOT `--body`. Plain text markdown will not render properly in email clients.\n\n## Template Structure\n\nThe agent should generate an HTML email body. Use inline styles (email clients strip `<style>` blocks).\n\n```html\n<div style=\"max-width:640px;margin:0 auto;font-family:-apple-system,BlinkMacSystemFont,'Segoe UI',Roboto,sans-serif;color:#1a1a1a;line-height:1.6\">\n\n  <h1 style=\"font-size:22px;border-bottom:2px solid #e5e5e5;padding-bottom:8px\">\n    🐉 {{TITLE}}\n  </h1>\n\n  <!-- Optional: Executive Summary for weekly -->\n  <p style=\"color:#555;font-size:14px;background:#f8f9fa;padding:12px;border-radius:6px\">\n    {{SUMMARY}}\n  </p>\n\n  <!-- Topic Section -->\n  <h2 style=\"font-size:17px;margin-top:24px;color:#333\">{{emoji}} {{label}}</h2>\n  <ul style=\"padding-left:20px\">\n    <li style=\"margin-bottom:10px\">\n      <strong>{{title}}</strong> — {{description}}\n      <br><a href=\"{{link}}\" style=\"color:#0969da;font-size:13px\">{{link}}</a>\n    </li>\n  </ul>\n\n  <!-- Repeat for each topic -->\n\n  <!-- KOL Section: Read metrics from twitter JSON data (metrics.impression_count, reply_","readmeExcerpt":"Skill: Tech News Digest Owner: dinstein Summary: Generate tech news digests with unified source model, quality scoring, and multi-format output. Five-layer data collection from RSS feeds, Twitter/X KOLs, Gi... Tags: latest:3.11.0 Version history: v3.11.0 | 2026-02-28T16:22:08.196Z | user Tavily backend, quality scores, domain limit fix, tests, CI v3.10.3 | 2026-02-27T15:46:08.147Z | user Docs alignment + config namin","codeSnippets":[],"executableExamples":[{"language":"bash","snippet":"mkdir -p workspace/config\n   cp config/defaults/sources.json workspace/config/tech-news-digest-sources.json\n   cp config/defaults/topics.json workspace/config/tech-news-digest-topics.json"},{"language":"bash","snippet":"# Unified pipeline (recommended) — runs all 5 sources in parallel + merge\n   python3 scripts/run-pipeline.py \\\n     --defaults config/defaults \\\n     --config workspace/config \\\n     --hours 48 --freshness pd \\\n     --archive-dir workspace/archive/tech-news-digest/ \\\n     --output /tmp/td-merged.json --verbose --force"},{"language":"json","snippet":"{\n  \"sources\": [\n    {\n      \"id\": \"openai-rss\",\n      \"type\": \"rss\",\n      \"name\": \"OpenAI Blog\",\n      \"url\": \"https://openai.com/blog/rss.xml\",\n      \"enabled\": true,\n      \"priority\": true,\n      \"topics\": [\"llm\", \"ai-agent\"],\n      \"note\": \"Official OpenAI updates\"\n    },\n    {\n      \"id\": \"sama-twitter\",\n      \"type\": \"twitter\", \n      \"name\": \"Sam Altman\",\n      \"handle\": \"sama\",\n      \"enabled\": true,\n      \"priority\": true,\n      \"topics\": [\"llm\", \"frontier-tech\"],\n      \"note\": \"OpenAI CEO\"\n    }\n  ]\n}"},{"language":"json","snippet":"{\n  \"topics\": [\n    {\n      \"id\": \"llm\",\n      \"emoji\": \"🧠\",\n      \"label\": \"LLM / Large Models\",\n      \"description\": \"Large Language Models, foundation models, breakthroughs\",\n      \"search\": {\n        \"queries\": [\"LLM latest news\", \"large language model breakthroughs\"],\n        \"must_include\": [\"LLM\", \"large language model\", \"foundation model\"],\n        \"exclude\": [\"tutorial\", \"beginner guide\"]\n      },\n      \"display\": {\n        \"max_items\": 8,\n        \"style\": \"detailed\"\n      }\n    }\n  ]\n}"},{"language":"bash","snippet":"python3 scripts/run-pipeline.py \\\n  --defaults config/defaults [--config CONFIG_DIR] \\\n  --hours 48 --freshness pd \\\n  --archive-dir workspace/archive/tech-news-digest/ \\\n  --output /tmp/td-merged.json --verbose --force"},{"language":"bash","snippet":"python3 scripts/fetch-rss.py [--defaults DIR] [--config DIR] [--hours 48] [--output FILE] [--verbose]"}],"parameters":null,"dependencies":[],"permissions":[],"extractedFiles":[{"path":"SKILL.md","content":"---\nname: tech-news-digest\ndescription: Generate tech news digests with unified source model, quality scoring, and multi-format output. Five-layer data collection from RSS feeds, Twitter/X KOLs, GitHub releases, Reddit, and web search. Pipeline-based scripts with retry mechanisms and deduplication. Supports Discord, email, and markdown templates.\nversion: \"3.11.0\"\nhomepage: https://github.com/draco-agent/tech-news-digest\nsource: https://github.com/draco-agent/tech-news-digest\nmetadata:\n  openclaw:\n    requires:\n      bins: [\"python3\"]\n    optionalBins: [\"mail\", \"msmtp\", \"gog\", \"gh\", \"openssl\", \"weasyprint\"]\nenv:\n  - name: TWITTER_API_BACKEND\n    required: false\n    description: \"Twitter API backend: 'official', 'twitterapiio', or 'auto' (default: auto)\"\n  - name: X_BEARER_TOKEN\n    required: false\n    description: Twitter/X API bearer token for KOL monitoring (official backend)\n  - name: TWITTERAPI_IO_KEY\n    required: false\n    description: twitterapi.io API key for KOL monitoring (twitterapiio backend)\n  - name: TAVILY_API_KEY\n    required: false\n    description: Tavily Search API key (alternative to Brave)\n  - name: WEB_SEARCH_BACKEND\n    required: false\n    description: \"Web search backend: auto (default), brave, or tavily\"\n  - name: BRAVE_API_KEYS\n    required: false\n    description: Brave Search API keys (comma-separated for rotation)\n  - name: BRAVE_API_KEY\n    required: false\n    description: Brave Search API key (single key fallback)\n  - name: GITHUB_TOKEN\n    required: false\n    description: GitHub token for higher API rate limits (auto-generated from GitHub App if not set)\n  - name: GH_APP_ID\n    required: false\n    description: GitHub App ID for automatic installation token generation\n  - name: GH_APP_INSTALL_ID\n    required: false\n    description: GitHub App Installation ID for automatic token generation\n  - name: GH_APP_KEY_FILE\n    required: false\n    description: Path to GitHub App private key PEM file\ntools:\n  - python3: Required. Runs data collection and merge scripts.\n  - mail: Optional. msmtp-based mail command for email delivery (preferred).\n  - gog: Optional. Gmail CLI for email delivery (fallback if mail not available).\nfiles:\n  read:\n    - config/defaults/: Default source and topic configurations\n    - references/: Prompt templates and output templates\n    - scripts/: Python pipeline scripts\n    - <workspace>/archive/tech-news-digest/: Previous digests for dedup\n  write:\n    - /tmp/td-*.json: Temporary pipeline intermediate outputs\n    - /tmp/td-email.html: Temporary email HTML body\n    - /tmp/td-digest.pdf: Generated PDF digest\n    - <workspace>/archive/tech-news-digest/: Saved digest archives\n---\n\n# Tech News Digest\n\nAutomated tech news digest system with unified data source model, quality scoring pipeline, and template-based output generation.\n\n## Quick Start\n\n1. **Configuration Setup**: Default configs are in `config/defaults/`. Copy to workspace for customization:\n   ```bash\n   mkdir -p workspace/config\n   cp config/d"},{"path":"README.md","content":"# Tech News Digest\n\n> Automated tech news digest — 138 sources, 5-layer pipeline, one chat message to install.\n\n**English** | [中文](README_CN.md)\n\n[![Tests](https://github.com/draco-agent/tech-news-digest/actions/workflows/test.yml/badge.svg)](https://github.com/draco-agent/tech-news-digest/actions/workflows/test.yml)\n[![Python 3.8+](https://img.shields.io/badge/python-3.8+-blue.svg)](https://www.python.org/downloads/)\n[![ClawHub](https://img.shields.io/badge/ClawHub-tech--news--digest-blueviolet)](https://clawhub.com/draco-agent/tech-news-digest)\n[![MIT License](https://img.shields.io/badge/license-MIT-green.svg)](LICENSE)\n\n## 💬 Install in One Message\n\nTell your [OpenClaw](https://openclaw.ai) AI assistant:\n\n> **\"Install tech-news-digest and send a daily digest to #tech-news every morning at 9am\"**\n\nThat's it. Your bot handles installation, configuration, scheduling, and delivery — all through conversation.\n\nMore examples:\n\n> 🗣️ \"Set up a weekly AI digest, only LLM and AI Agent topics, deliver to Discord #ai-weekly every Monday\"\n\n> 🗣️ \"Install tech-news-digest, add my RSS feeds, and send crypto news to Telegram\"\n\n> 🗣️ \"Give me a tech digest right now, skip Twitter sources\"\n\nOr install via CLI:\n```bash\nclawhub install tech-news-digest\n```\n\n## 📊 What You Get\n\nA quality-scored, deduplicated tech digest built from **138 sources**:\n\n| Layer | Sources | What |\n|-------|---------|------|\n| 📡 RSS | 49 feeds | OpenAI, Anthropic, Ben's Bites, HN, 36氪, CoinDesk… |\n| 🐦 Twitter/X | 48 KOLs | @karpathy, @VitalikButerin, @sama, @elonmusk… |\n| 🔍 Web Search | 4 topics | Brave Search API with freshness filters |\n| 🐙 GitHub | 28 repos | Releases from key projects (LangChain, vLLM, DeepSeek, Llama…) |\n| 🗣️ Reddit | 13 subs | r/MachineLearning, r/LocalLLaMA, r/CryptoCurrency… |\n\n### Pipeline\n\n```\n       run-pipeline.py (~30s)\n              ↓\n  RSS ─┐\n  Twitter ─┤\n  Web ─────┤── parallel fetch ──→ merge-sources.py\n  GitHub ──┤\n  Reddit ──┘\n              ↓\n  Quality Scoring → Deduplication → Topic Grouping\n              ↓\n    Discord / Email / PDF output\n```\n\n**Quality scoring**: priority source (+3), multi-source cross-ref (+5), recency (+2), engagement (+1), Reddit score bonus (+1/+3/+5), already reported (-5).\n\n## ⚙️ Configuration\n\n- `config/defaults/sources.json` — 138 built-in sources\n- `config/defaults/topics.json` — 4 topics with search queries & Twitter queries\n- User overrides in `workspace/config/` take priority\n\n## 🎨 Customize Your Sources\n\nWorks out of the box with 138 built-in sources — but fully customizable. Copy the defaults to your workspace config and override:\n\n```bash\n# Copy and customize\ncp config/defaults/sources.json workspace/config/tech-news-digest-sources.json\ncp config/defaults/topics.json workspace/config/tech-news-digest-topics.json\n```\n\nYour overlay file **merges** with defaults:\n- **Override** a source by matching its `id` — your version replaces the default\n- **Add** new sources with a unique `id` — appended to the list\n- **Di"},{"path":"_meta.json","content":"{\n  \"ownerId\": \"kn74589cx1nbhnc3x0f3nwre39814699\",\n  \"slug\": \"tech-news-digest\",\n  \"version\": \"3.11.0\",\n  \"publishedAt\": 1772295728196\n}"},{"path":"references/digest-prompt.md","content":"# Digest Prompt Template\n\nReplace `<...>` placeholders before use. Daily defaults shown; weekly overrides in parentheses.\n\n## Placeholders\n\n| Placeholder | Default | Weekly Override |\n|-------------|---------|----------------|\n| `<MODE>` | `daily` | `weekly` |\n| `<TIME_WINDOW>` | `past 1-2 days` | `past 7 days` |\n| `<FRESHNESS>` | `pd` | `pw` |\n| `<RSS_HOURS>` | `48` | `168` |\n| `<ITEMS_PER_SECTION>` | `3-5` | `5-8` |\n| `<BLOG_PICKS_COUNT>` | `2-3` | `3-5` |\n| `<EXTRA_SECTIONS>` | *(none)* | `📊 Weekly Trend Summary` |\n| `<SUBJECT>` | `Daily Tech Digest - YYYY-MM-DD` | `Weekly Tech Digest - YYYY-MM-DD` |\n| `<WORKSPACE>` | Your workspace path | |\n| `<SKILL_DIR>` | Installed skill directory | |\n| `<DISCORD_CHANNEL_ID>` | Target channel ID | |\n| `<EMAIL>` | *(optional)* Recipient email | |\n| `<EMAIL_FROM>` | *(optional)* e.g. `MyBot <bot@example.com>` | |\n| `<LANGUAGE>` | `Chinese` | |\n| `<TEMPLATE>` | `discord` / `email` / `markdown` | |\n| `<DATE>` | Today's date YYYY-MM-DD (caller provides) | |\n| `<VERSION>` | Read from SKILL.md frontmatter | |\n\n---\n\nGenerate the <MODE> tech digest for **<DATE>**. Use `<DATE>` as the report date — do NOT infer it.\n\n## Configuration\n\nRead config files (workspace overrides take priority over defaults):\n1. **Sources**: `<WORKSPACE>/config/tech-news-digest-sources.json` → fallback `<SKILL_DIR>/config/defaults/sources.json`\n2. **Topics**: `<WORKSPACE>/config/tech-news-digest-topics.json` → fallback `<SKILL_DIR>/config/defaults/topics.json`\n\n## Context: Previous Report\n\nRead the most recent file from `<WORKSPACE>/archive/tech-news-digest/` to avoid repeats and follow up on developing stories. Skip if none exists.\n\n## Data Collection Pipeline\n\n**Use the unified pipeline** (runs all 5 sources in parallel, ~30s):\n\n```bash\npython3 <SKILL_DIR>/scripts/run-pipeline.py \\\n  --defaults <SKILL_DIR>/config/defaults \\\n  --config <WORKSPACE>/config \\\n  --hours <RSS_HOURS> --freshness <FRESHNESS> \\\n  --archive-dir <WORKSPACE>/archive/tech-news-digest/ \\\n  --output /tmp/td-merged.json --verbose --force\n```\n\nIf it fails, run individual scripts in `<SKILL_DIR>/scripts/` (see each script's `--help`), then merge with `merge-sources.py`.\n\n## Report Generation\n\nGet a structured overview:\n```bash\npython3 <SKILL_DIR>/scripts/summarize-merged.py --input /tmp/td-merged.json --top <ITEMS_PER_SECTION>\n```\n\nUse this output to select articles — **do NOT write ad-hoc Python to parse the JSON**. Apply the template from `<SKILL_DIR>/references/templates/<TEMPLATE>.md`.\n\nSelect articles **purely by quality_score regardless of source type**. Articles in merged JSON are already sorted by quality_score descending within each topic — respect this order. For Reddit posts, append `*[Reddit r/xxx, {{score}}↑]*`.\n\nEach article line must include its quality score using 🔥 prefix. Format: `🔥{score} | {summary with link}`. This makes scoring transparent and helps readers identify the most important news at a glance.\n\n### Executive Summary\n2-4 sentences between t"},{"path":"references/templates/discord.md","content":"# Tech Digest Discord Template\n\nDiscord-optimized format with bullet points and link suppression.\n\n## Template Structure\n\n```markdown\n# 🚀 Tech Digest - {{DATE}}\n\n{{#topics}}\n## {{emoji}} {{label}}\n\n{{#articles}}\n• 🔥{{quality_score}} | {{title}}\n  <{{link}}>\n  {{#multi_source}}*[{{source_count}} sources]*{{/multi_source}}\n\n{{/articles}}\n{{/topics}}\n\n---\n📊 Data Sources: RSS {{rss_count}} | Twitter {{twitter_count}} | Reddit {{reddit_count}} | Web {{web_count}} | GitHub {{github_count}} releases | After dedup: {{merged_count}} articles\n🤖 Generated by tech-news-digest v{{version}} | <https://github.com/draco-agent/tech-news-digest> | Powered by OpenClaw\n```\n\n## Delivery\n\n- **Default: Channel** — Send to the Discord channel specified by `DISCORD_CHANNEL_ID`\n- Use `message` tool with `target` set to the channel ID for channel delivery\n- For DM delivery instead, set `target` to a user ID\n\n## Discord-Specific Features\n\n- **Link suppression**: Wrap links in `<>` to prevent embeds\n- **Bullet format**: Use `•` for clean mobile display  \n- **No tables**: Discord mobile doesn't handle markdown tables well\n- **Emoji headers**: Visual hierarchy with topic emojis\n- **Concise metadata**: Source count and multi-source indicators\n- **Character limits**: Discord messages have 2000 char limit, may need splitting\n\n## Example Output\n\n```markdown\n# 🚀 Tech Digest - 2026-02-15\n\n## 🧠 LLM / Large Models\n\n• 🔥15 | OpenAI releases GPT-5 with breakthrough reasoning capabilities\n  <https://openai.com/blog/gpt5-announcement>\n  *[3 sources]*\n\n• 🔥12 | Meta's Llama 3.1 achieves new MMLU benchmarks\n  <https://ai.meta.com/blog/llama-31-release>\n\n## 🤖 AI Agent\n\n• 🔥14 | LangChain launches production-ready agent framework\n  <https://blog.langchain.dev/production-agents>\n\n## 💰 Cryptocurrency\n\n• 🔥18 | Bitcoin reaches new ATH at $67,000 amid ETF approval\n  <https://coindesk.com/markets/btc-ath-etf>\n  *[2 sources]*\n\n## 📢 KOL Updates\n\n• **Elon Musk** (@elonmusk) — Confirmed X's crypto trading feature `👁 2.1M | 💬 12.3K | 🔁 8.5K | ❤️ 49.8K`\n  <https://twitter.com/elonmusk/status/123456789>\n• **@saylor** — Valentine's BTC enthusiasm `👁 450K | 💬 1.2K | 🔁 3.1K | ❤️ 13K`\n  <https://twitter.com/saylor/status/987654321>\n\n---\n📊 Data Sources: RSS 285 | Twitter 67 | Reddit 45 | Web 60 | GitHub 29 releases | After dedup: 95 articles\n```\n\n## Variables\n\n- `{{DATE}}` - Report date (YYYY-MM-DD format)\n- `{{topics}}` - Array of topic objects\n- `{{emoji}}` - Topic emoji \n- `{{label}}` - Topic display name\n- `{{articles}}` - Array of article objects per topic\n- `{{title}}` - Article title (truncated if needed)\n- `{{link}}` - Article URL\n- `{{quality_score}}` - Article quality score (higher = more important)\n- `{{multi_source}}` - Boolean, true if article from multiple sources\n- `{{source_count}}` - Number of sources for this article\n- `{{total_sources}}` - Total number of sources used\n- `{{total_articles}}` - Total articles in digest"}],"languages":[],"docsSourceLabel":"CLAWHUB","editorialOverview":"Generate tech news digests with unified source model, quality scoring, and multi-format output. Five-layer data collection from RSS feeds, Twitter/X KOLs, Gi... Skill: Tech News Digest Owner: dinstein Summary: Generate tech news digests with unified source model, quality scoring, and multi-format output. Five-layer data collection from RSS feeds, Twitter/X KOLs, Gi... Tags: latest:3.11.0 Version history: v3.11.0 | 2026-02-28T16:22:08.196Z | user Tavily backend, quality scores, domain limit fix, tests, CI v3.10.3 | 2026-02-27T15:46:08.147Z | user Docs alignment + config namin","editorialQuality":{"score":100,"threshold":65,"status":"ready","wordCount":1222,"uniquenessScore":47,"reasons":[]}},"media":{"evidence":{"source":"no-media","verified":false,"confidence":"low","updatedAt":"2026-04-15T00:45:39.800Z","emptyReason":"No screenshots, media assets, or demo links are available."},"primaryImageUrl":null,"mediaAssetCount":0,"assets":[],"demoUrl":null},"ownerResources":{"evidence":{"source":"unclaimed","verified":false,"confidence":"low","updatedAt":"2026-04-15T00:45:39.800Z","emptyReason":"This page has not been claimed by the agent owner."},"hasCustomPage":false,"customPageUpdatedAt":null,"customLinks":[],"structuredLinks":{"docsUrl":null,"demoUrl":null,"supportUrl":null,"pricingUrl":null,"statusUrl":null},"customPage":null},"relatedAgents":{"evidence":{"source":"protocol-neighbors","verified":false,"confidence":"medium","updatedAt":"2026-10-09T14:02:08.759Z","emptyReason":null},"items":[{"id":"b917f68a-ebff-438e-84f8-3f4b2494c0bc","entityType":"agent","canonicalPath":"/agent/activepieces-activepieces","slug":"activepieces-activepieces","name":"activepieces","description":"AI Agents & MCPs & AI Workflow Automation • (~400 MCP servers for AI agents) • AI Automation / AI Agent with MCPs • AI Workflows & AI Agents • MCPs for AI Agents","url":"https://github.com/activepieces/activepieces","homepage":"https://www.activepieces.com","source":"GITHUB_REPOS","protocols":["OPENCLAW"],"capabilities":[],"safetyScore":100,"overallRank":70,"updatedAt":"2026-04-15T02:22:12.426Z","createdAt":"2026-02-25T03:38:12.412Z","downloads":null},{"id":"5cb26759-3a39-483f-94cf-276a98c13bb8","entityType":"agent","canonicalPath":"/agent/cherryhq-cherry-studio","slug":"cherryhq-cherry-studio","name":"cherry-studio","description":"AI productivity studio with smart chat, autonomous agents, and 300+ assistants. Unified access to frontier LLMs","url":"https://github.com/CherryHQ/cherry-studio","homepage":"https://cherry-ai.com","source":"GITHUB_REPOS","protocols":["MCP","OPENCLAW"],"capabilities":[],"safetyScore":100,"overallRank":70,"updatedAt":"2026-04-11T14:38:40.986Z","createdAt":"2026-02-25T03:38:19.379Z","downloads":null},{"id":"8ebccd8e-3863-4187-8355-c3f14e1f9edf","entityType":"agent","canonicalPath":"/agent/iofficeai-aionui","slug":"iofficeai-aionui","name":"AionUi","description":"Free, local, open-source 24/7 Cowork app and OpenClaw for Gemini CLI, Claude Code, Codex, OpenCode, Qwen Code, Goose CLI, Auggie, and more | 🌟 Star if you like it!","url":"https://github.com/iOfficeAI/AionUi","homepage":"https://www.aionui.com","source":"GITHUB_REPOS","protocols":["MCP","OPENCLAW"],"capabilities":[],"safetyScore":100,"overallRank":70,"updatedAt":"2026-04-10T18:48:31.762Z","createdAt":"2026-02-25T03:38:16.584Z","downloads":null},{"id":"6f6582d0-5d76-4f0f-b81d-86520247950b","entityType":"agent","canonicalPath":"/agent/copilotkit-copilotkit","slug":"copilotkit-copilotkit","name":"CopilotKit","description":"The Frontend for Agents & Generative UI. React + Angular","url":"https://github.com/CopilotKit/CopilotKit","homepage":"https://docs.copilotkit.ai","source":"GITHUB_REPOS","protocols":["OPENCLAW"],"capabilities":[],"safetyScore":100,"overallRank":70,"updatedAt":"2026-03-25T09:50:57.846Z","createdAt":"2026-02-25T03:39:14.617Z","downloads":null}],"links":{"hub":"/agent","source":"/agent/source/clawhub","protocols":[{"label":"OpenClaw","href":"/agent/protocol/openclew"}]}}}