{"id":"5478eb6e-8e19-4de1-b15c-ce8d694584cd","entityType":"agent","slug":"clawhub-pruna-ai-stable-audio-2-5","name":"stable-audio-2.5","canonicalUrl":"https://www.xpersona.co/agent/clawhub-pruna-ai-stable-audio-2-5","canonicalPath":"/agent/clawhub-pruna-ai-stable-audio-2-5","generatedAt":"2026-10-10T14:46:32.564Z","source":"CLAWHUB","claimStatus":"UNCLAIMED","verificationTier":"NONE","summary":{"evidence":{"source":"editorial-content","verified":true,"confidence":"high","updatedAt":"2026-10-10T12:32:11.870Z","emptyReason":null},"description":"Use when someone wants light instrumental background music — an ambient bed under dialogue or underscore for reels and explainers. Skill: stable-audio-2.5 Owner: pruna-ai Summary: Use when someone wants light instrumental background music — an ambient bed under dialogue or underscore for reels and explainers. Tags: ai:1.0.14, generative:1.0.14, latest:1.0.14, pruna:1.0.14 Version history: v1.0.14 | 2026-09-29T15:36:14.516Z | auto - Updated version metadata from 1.0.13 to 1.0.14 in SKILL.md. - Removed the file skill-card.md. v1.0.13 | 2026-09-17T","descriptionLabel":"Technical summary","evidenceSummary":"Capability contract not published. No trust telemetry is available yet. 1.4K downloads reported by the source. Last updated 10/10/2026.","installCommand":"clawhub skill install s173djpkkm1x4yfzg522h5vfp989mrn5:stable-audio-2-5","sourceUrl":"https://clawhub.ai/pruna-ai/stable-audio-2-5","homepage":"https://clawhub.ai/pruna-ai/skills/stable-audio-2-5","primaryLinks":[{"label":"View on ClawHub","url":"https://clawhub.ai/pruna-ai/stable-audio-2-5","kind":"source"},{"label":"Homepage","url":"https://clawhub.ai/pruna-ai/skills/stable-audio-2-5","kind":"homepage"}],"safetyScore":84,"overallRank":62,"popularityScore":63,"trustScore":null,"claimedByName":null,"isOwner":false,"seoDescription":"Use when someone wants light instrumental background music — an ambient bed under dialogue or underscore for reels and explainers. Skill: stable-audio-2.5 Owner"},"coverage":{"evidence":{"source":"public-profile","verified":false,"confidence":"medium","updatedAt":"2026-10-10T12:32:11.870Z","emptyReason":null},"protocols":[{"protocol":"OPENCLEW","label":"OpenClaw","status":"self-declared","notes":"Declared in the public agent profile."}],"capabilities":[],"verifiedCount":0,"selfDeclaredCount":1,"capabilityMatrix":{"rows":[{"key":"OPENCLEW","type":"protocol","support":"unknown","confidenceSource":"profile","notes":"Listed on profile"}],"flattenedTokens":"protocol:OPENCLEW|unknown|profile"}},"adoption":{"evidence":{"source":"CLAWHUB","verified":false,"confidence":"medium","updatedAt":"2026-10-10T12:32:11.870Z","emptyReason":null},"stars":null,"forks":null,"downloads":1431,"packageName":null,"latestVersion":"1.0.14","tractionLabel":"1.4K downloads"},"release":{"evidence":{"source":"CLAWHUB","verified":false,"confidence":"medium","updatedAt":"2026-10-10T12:32:11.797Z","emptyReason":null},"lastUpdatedAt":"2026-10-10T12:32:11.870Z","lastCrawledAt":"2026-10-10T12:32:11.797Z","lastIndexedAt":null,"nextCrawlAt":"2026-10-11T12:32:11.797Z","lastVerifiedAt":null,"highlights":[{"version":"1.0.14","createdAt":"2026-09-29T15:36:14.516Z","changelog":"- Updated version metadata from 1.0.13 to 1.0.14 in SKILL.md. - Removed the file skill-card.md.","fileCount":4,"zipByteSize":3540},{"version":"1.0.13","createdAt":"2026-09-17T13:59:09.590Z","changelog":"Stable-audio-2-5 v1.0.13 - Version updated to 1.0.13 in SKILL.md metadata. - Removed redundant skill-card.md file. - No changes to usage, features, or user workflow.","fileCount":4,"zipByteSize":3673},{"version":"1.0.12","createdAt":"2026-09-10T13:57:42.529Z","changelog":"stable-audio-2-5 v1.0.12 - Updated version number to 1.0.12 in SKILL.md. - Removed the sample file skill-card.md from the repository. - No changes to usage, features, or functionality.","fileCount":4,"zipByteSize":3673},{"version":"1.0.11","createdAt":"2026-09-03T14:11:57.797Z","changelog":"stable-audio-2-5 v1.0.11 changelog: - Updated version to 1.0.11 in metadata. - Removed redundant skill-card.md file, consolidating documentation into SKILL.md.","fileCount":4,"zipByteSize":3727},{"version":"1.0.10","createdAt":"2026-08-28T07:57:47.430Z","changelog":"- Bumped version to 1.0.10 in metadata. - Removed the file: skill-card.md. - No changes to features or functionality.","fileCount":4,"zipByteSize":3702},{"version":"1.0.9","createdAt":"2026-08-04T06:18:29.963Z","changelog":"- Update version to 1.0.9 in metadata. - Remove skill-card.md file from the repository.","fileCount":4,"zipByteSize":3772},{"version":"1.0.8","createdAt":"2026-07-28T17:20:34.327Z","changelog":"- Added explicit instruction to open a generation-diversity clarification intake before the first POST request. - Updated agent habit guidance to initiate clarification intake with generation-diversity before generation. - Removed the skill-card.md file, reducing redundant metadata. - Updated version number to 1.0.8.","fileCount":4,"zipByteSize":3703},{"version":"1.0.7","createdAt":"2026-07-23T12:35:36.093Z","changelog":"stable-audio-2-5 v1.0.7 is a major update streamlining documentation, clarifying agent prompts, and centralizing generation policies. - Simplified the SKILL.md with direct prerequisites, concise input requirements, and explicit When NOT to use section. - Replaced embedded policy descriptions with references to dedicated skills (generation-diversity, audio-prompting, pruna-api) to reduce duplication. - Removed 8 docs (references, policy guides, and skill-card) now replaced by skill references or external links. - Updated environment setup and guidance for first agent reply. - Added links and installation commands for related skills and next steps.","fileCount":4,"zipByteSize":3730}]},"execution":{"evidence":{"source":"CLAWHUB","verified":false,"confidence":"low","updatedAt":null,"emptyReason":"No published capability contract is available yet."},"installCommand":"clawhub skill install s173djpkkm1x4yfzg522h5vfp989mrn5:stable-audio-2-5","setupComplexity":"low","setupSteps":["Setup complexity is classified as HIGH. You must provision dedicated cloud infrastructure or an isolated VM. Do not run this directly on your local workstation.","Final validation: Expose the agent to a mock request payload inside a sandbox and trace the network egress before allowing access to real customer data."],"contract":{"contractStatus":"missing","authModes":[],"requires":[],"forbidden":[],"supportsMcp":false,"supportsA2a":false,"supportsStreaming":false,"inputSchemaRef":null,"outputSchemaRef":null,"dataRegion":null,"contractUpdatedAt":null,"sourceUpdatedAt":null,"freshnessSeconds":null},"invocationGuide":{"preferredApi":{"snapshotUrl":"https://www.xpersona.co/api/v1/agents/clawhub-pruna-ai-stable-audio-2-5/snapshot","contractUrl":"https://www.xpersona.co/api/v1/agents/clawhub-pruna-ai-stable-audio-2-5/contract","trustUrl":"https://www.xpersona.co/api/v1/agents/clawhub-pruna-ai-stable-audio-2-5/trust"},"curlExamples":["curl -s \"https://www.xpersona.co/api/v1/agents/clawhub-pruna-ai-stable-audio-2-5/snapshot\"","curl -s \"https://www.xpersona.co/api/v1/agents/clawhub-pruna-ai-stable-audio-2-5/contract\"","curl -s \"https://www.xpersona.co/api/v1/agents/clawhub-pruna-ai-stable-audio-2-5/trust\""],"jsonRequestTemplate":{"query":"summarize this repo","constraints":{"maxLatencyMs":2000,"protocolPreference":["OPENCLEW"]}},"jsonResponseTemplate":{"ok":true,"result":{"summary":"...","confidence":0.9},"meta":{"source":"CLAWHUB","generatedAt":"2026-10-10T14:46:32.560Z"}},"retryPolicy":{"maxAttempts":3,"backoffMs":[500,1500,3500],"retryableConditions":["HTTP_429","HTTP_503","NETWORK_TIMEOUT"]}},"endpoints":{"dossierUrl":"https://www.xpersona.co/api/v1/agents/clawhub-pruna-ai-stable-audio-2-5/dossier","snapshotUrl":"https://www.xpersona.co/api/v1/agents/clawhub-pruna-ai-stable-audio-2-5/snapshot","contractUrl":"https://www.xpersona.co/api/v1/agents/clawhub-pruna-ai-stable-audio-2-5/contract","trustUrl":"https://www.xpersona.co/api/v1/agents/clawhub-pruna-ai-stable-audio-2-5/trust"}},"reliability":{"evidence":{"source":"runtime-metrics","verified":false,"confidence":"low","updatedAt":null,"emptyReason":"No trust, reliability, or runtime telemetry is available."},"trust":{"status":"unavailable","handshakeStatus":"UNKNOWN","verificationFreshnessHours":null,"reputationScore":null,"p95LatencyMs":null,"successRate30d":null,"fallbackRate":null,"attempts30d":null,"trustUpdatedAt":null,"trustConfidence":"unknown","sourceUpdatedAt":null,"freshnessSeconds":null},"decisionGuardrails":{"doNotUseIf":["Contract metadata is missing or unavailable for deterministic execution."],"safeUseWhen":[],"riskFlags":["missing_or_unavailable_contract","trust_data_unavailable","schema_references_missing"],"operationalConfidence":"low"},"executionMetrics":{"observedLatencyMsP50":null,"observedLatencyMsP95":null,"estimatedCostUsd":null,"uptime30d":null,"rateLimitRpm":null,"rateLimitBurst":null,"lastVerifiedAt":null,"verificationSource":null},"runtimeMetrics":{"successRate":null,"avgLatencyMs":null,"avgCostUsd":null,"hallucinationRate":null,"retryRate":null,"disputeRate":null,"p50Latency":null,"p95Latency":null,"lastUpdated":null}},"benchmarks":{"evidence":{"source":"no-benchmark-data","verified":false,"confidence":"low","updatedAt":null,"emptyReason":"No benchmark suites or observed failure patterns are available."},"suites":[],"failurePatterns":[]},"artifacts":{"evidence":{"source":"CLAWHUB","verified":false,"confidence":"high","updatedAt":"2026-10-10T12:32:11.870Z","emptyReason":null},"readme":"Skill: stable-audio-2.5\n\nOwner: pruna-ai\n\nSummary: Use when someone wants light instrumental background music — an ambient bed under dialogue or underscore for reels and explainers.\n\nTags: ai:1.0.14, generative:1.0.14, latest:1.0.14, pruna:1.0.14\n\nVersion history:\n\nv1.0.14 | 2026-09-29T15:36:14.516Z | auto\n\n- Updated version metadata from 1.0.13 to 1.0.14 in SKILL.md.\n- Removed the file skill-card.md.\n\nv1.0.13 | 2026-09-17T13:59:09.590Z | auto\n\nStable-audio-2-5 v1.0.13\n\n- Version updated to 1.0.13 in SKILL.md metadata.\n- Removed redundant skill-card.md file.\n- No changes to usage, features, or user workflow.\n\nv1.0.12 | 2026-09-10T13:57:42.529Z | auto\n\nstable-audio-2-5 v1.0.12\n\n- Updated version number to 1.0.12 in SKILL.md.\n- Removed the sample file skill-card.md from the repository.\n- No changes to usage, features, or functionality.\n\nv1.0.11 | 2026-09-03T14:11:57.797Z | auto\n\nstable-audio-2-5 v1.0.11 changelog:\n\n- Updated version to 1.0.11 in metadata.\n- Removed redundant skill-card.md file, consolidating documentation into SKILL.md.\n\nv1.0.10 | 2026-08-28T07:57:47.430Z | auto\n\n- Bumped version to 1.0.10 in metadata.\n- Removed the file: skill-card.md.\n- No changes to features or functionality.\n\nv1.0.9 | 2026-08-04T06:18:29.963Z | auto\n\n- Update version to 1.0.9 in metadata.\n- Remove skill-card.md file from the repository.\n\nv1.0.8 | 2026-07-28T17:20:34.327Z | auto\n\n- Added explicit instruction to open a generation-diversity clarification intake before the first POST request.\n- Updated agent habit guidance to initiate clarification intake with generation-diversity before generation.\n- Removed the skill-card.md file, reducing redundant metadata. \n- Updated version number to 1.0.8.\n\nv1.0.7 | 2026-07-23T12:35:36.093Z | auto\n\nstable-audio-2-5 v1.0.7 is a major update streamlining documentation, clarifying agent prompts, and centralizing generation policies.\n\n- Simplified the SKILL.md with direct prerequisites, concise input requirements, and explicit When NOT to use section.\n- Replaced embedded policy descriptions with references to dedicated skills (generation-diversity, audio-prompting, pruna-api) to reduce duplication.\n- Removed 8 docs (references, policy guides, and skill-card) now replaced by skill references or external links.\n- Updated environment setup and guidance for first agent reply.\n- Added links and installation commands for related skills and next steps.\n\nv1.0.6 | 2026-07-16T21:00:18.681Z | auto\n\n- Added a shared generation policy section detailing required steps before paid predictions: random seed ritual, diversity rotation, and quality checklist.\n- Introduced new reference docs: generation diversity, quality checklists, and random seed ritual.\n- Updated related documentation links to reflect clearer organization and shared policy usage.\n- Removed redundant or outdated files, including skill-card.md.\n- Clarified language around when and how to use this skill for instrumental background beds.\n\nv1.0.2 | 2026-07-16T13:35:16.613Z | auto\n\n- Bumped version to 1.0.2.\n- Documentation updates in SKILL.md and references.\n- Removed the redundant skill-card.md file for improved clarity.\n\nv1.0.1 | 2026-07-14T15:49:41.469Z | auto\n\n- Update to version 1.0.1 with improved documentation and reference links\n- Replaced local script and skill references with direct GitHub links for easier navigation\n- Removed legacy and duplicate files (pspm.json, skill-card.md)\n- Updated related skill and workflow sections to reflect accurate current locations\n\nv0.0.1 | 2026-07-14T15:06:24.130Z | auto\n\n- Initial release of Stable Audio 2.5 skill for generating instrumental background music beds.\n- Supports text-to-music for light, ambient, or underscore-style audio—ideal for launch reels, explainers, and under dialogue/voiceover.\n- Integrates with Replicate’s stability-ai/stable-audio-2.5 model; outputs a single MP3.\n- Includes mix helper script to automatically generate and mix beds with ffmpeg.\n- Provides guide for prompts, integration instructions, and practical launch reel workflows.\n\nv1.0.0 | 2026-07-14T10:16:23.541Z | auto\n\n- Initial release of stable-audio-2.5 skill for generating instrumental background beds using the Stability AI/Replicate model.\n- Designed for adding light, ambient music under launch reels, dialogue, or explainers—focus on understated, vocal-free arrangements.\n- Includes usage instructions, environment setup, model input options, and HTTP (curl) API example.\n- Provides integration details for mixing music under video using `launch_background_music.py` and sample automation plans.\n- Prompt guidelines and related tool references included for best results.\n\nArchive index:\n\nArchive v1.0.14: 4 files, 3540 bytes\n\nFiles: skill-card.md (1648b), skill.manifest.json (23b), SKILL.md (4613b), _meta.json (136b)\n\nFile v1.0.14:SKILL.md\n\n---\nname: stable-audio-2.5\ndescription: Use when someone wants light instrumental background music — an ambient bed under dialogue or underscore for reels and explainers.\nlicense: MIT\nmetadata:\n  version: \"1.0.14\"\n  package: pruna-skills\n  provider: replicate\n  replicate_model: stability-ai/stable-audio-2.5\n---\n\n## Prerequisites\n\nInstall and load these skills before generating (skip if already in context via `@pruna`):\n\n| Skill | Description | Install |\n| --- | --- | --- |\n| `generation-diversity` | Use when writing any generative prompt — ritual seed, explicit structure, scenario axes, and quality gates before paid API calls. | `npx skills add PrunaAI/pruna-skills@generation-diversity -y` |\n| `audio-prompting` | Use when crafting TTS, music, or bed prompts for any generative audio model — director style, song structure, and post-production layering. | `npx skills add PrunaAI/pruna-skills@audio-prompting -y` |\n| `pruna-api` | Use before any Pruna or Replicate HTTP call — credentials, upload/poll/download, parallel batches, and agent safety. | `npx skills add PrunaAI/pruna-skills@pruna-api -y` |\n\nOr install the full suite once: `npx skills add PrunaAI/pruna-skills@pruna -y`\n\nFollow each skill's **Before generating** / craft sections — do not restate guide content here.\n\n## Agent habit\n\nIn the **first reply**, name `` `stable-audio-2.5` `` in backticks, confirm `REPLICATE_API_TOKEN` (or stop with signup links from `pruna-api`), then ask for required inputs. Open intake → **`generation-diversity`** clarification intake before the first `POST`. Redirect when **When NOT to use** fits better.\n\n## When NOT to use\n\nUse a different skill instead:\n\n| Skill | Description | Install |\n| --- | --- | --- |\n| `music-2.5` | Use when someone wants an original AI song with vocals — sung lyrics, a style prompt track, or source audio for a music video. | `npx skills add PrunaAI/pruna-skills@music-2.5 -y` |\n| `gemini-3.1-flash-tts` | Use when someone needs spoken narration or voiceover — explainer tracks, documentary lines, or voice to pair with generated video. | `npx skills add PrunaAI/pruna-skills@gemini-3.1-flash-tts -y` |\n\n## Environment\n\n```bash\nexport REPLICATE_API_TOKEN=r8_...\n```\n\nRequires **`ffmpeg`** and **`ffprobe`** on PATH for the mix step.\n\n## HTTP (curl)\n\n```bash\ncurl -s -X POST \\\n  -H \"Authorization: Bearer ${REPLICATE_API_TOKEN}\" \\\n  -H \"Content-Type: application/json\" \\\n  -d '{\n    \"input\": {\n      \"prompt\": \"Instrumental light electronic pop bed, soft groove and mellow synth pads, calm positive tech atmosphere, understated background music, no vocals, 94 BPM\",\n      \"duration\": 90,\n      \"steps\": 8,\n      \"cfg_scale\": 1\n    }\n  }' \\\n  \"https://api.replicate.com/v1/models/stability-ai/stable-audio-2.5/predictions\"\n```\n\nPoll `urls.get` until `status` is `succeeded`; download `output` MP3.\n\n## Before generating\n\n1. Complete Prerequisites guide reading order — bed prompt craft: `audio-prompting` **Worked examples** (instrumental bed).\n2. Confirm **`prompt`**, **`duration`** (match or slightly exceed reel length), and mix **`volume`** (~0.08–0.15 under VO). When listing fields, name **`REPLICATE_API_TOKEN`** (Replicate — not `PRUNA_API_KEY`).\n3. **Model notes:** lead with **Instrumental** and **no vocals**. Duration 1–190s. Prefer understated beds (BPM ~88–98 for tech launch reels) so music does not compete with dialogue.\n\n## Required input\n\n- `prompt` (string)\n\n## Common optional fields\n\n- `duration` — seconds, 1–190\n- `steps` — 4–8 (default 8)\n- `cfg_scale` — 1–25 (default 1)\n- `seed` — optional integer\n\n## Typical next steps\n\nCommon follow-ons after this skill:\n\n| Skill | Description | Install |\n| --- | --- | --- |\n| `gemini-3.1-flash-tts` | Use when someone needs spoken narration or voiceover — explainer tracks, documentary lines, or voice to pair with generated video. | `npx skills add PrunaAI/pruna-skills@gemini-3.1-flash-tts -y` |\n| `visual-transition-reel` | Use when someone wants a montage with transitions between shots — action-sequence reel or multi-scene piece where narration is optional. | `npx skills add PrunaAI/pruna-skills@visual-transition-reel -y` |\n| `video-editing` | Use when assembling or polishing already-rendered clips with ffmpeg — concat, crossfades, burned captions and subtitles, text/logo overlays, before/after sliders, background music beds, platform export — or when composing a multi-layer HTML combination video with Hyperframes. Not for AI video generation, prompt craft, or model-based video edits. | `npx skills add PrunaAI/pruna-skills@video-editing -y` |\n\nFile v1.0.14:_meta.json\n\n{\n  \"ownerId\": \"kn7cagwf7q3t0cxrgteb7xk0bh81j0eb\",\n  \"slug\": \"stable-audio-2-5\",\n  \"version\": \"1.0.14\",\n  \"publishedAt\": 1790696174516\n}\n\nFile v1.0.14:skill-card.md\n\n## Description:\n\nHelps create light instrumental background music for dialogue, reels, and explainers using Stable Audio 2.5.\n\nThis skill is ready for commercial/non-commercial use.\n\n## Publisher:\n\n[pruna-ai](https://clawhub.ai/user/pruna-ai)\n\n### License/Terms of Use:\n\nMIT-0\n\n## Use Case:\n\nCreators and editors use this skill to prepare instrumental music prompts and generation steps for unobtrusive beds beneath narration or short videos.\n\n### Deployment Geography for Use:\n\nGlobal\n\n## Known Risks and Mitigations:\n\nRisk: Unpinned skill installation commands may fetch changed dependencies.\n\nMitigation: Review the installation commands and use a trusted or pinned installation path when available.\n\nRisk: Prompts or credentials may be exposed when calling an external audio service.\n\nMitigation: Provide REPLICATE_API_TOKEN only when using Replicate, and omit secrets and private information from prompts.\n\n## Reference(s):\n\n- [ClawHub skill release](https://clawhub.ai/pruna-ai/skills/stable-audio-2-5)\n\n## Skill Output:\n\n**Output Type(s):** [Text, Shell commands, Configuration guidance]\n\n**Output Format:** [Markdown with bash and JSON examples]\n\n**Output Parameters:** [1D]\n\n**Other Properties Related to Output:** [Guides Replicate MP3 generation; optional mixing requires ffmpeg and ffprobe.]\n\n## Skill Version(s):\n\n1.0.14 (source: skill frontmatter and server release)\n\n## Ethical Considerations:\n\nUsers should evaluate whether this skill is appropriate for their environment, review any generated or modified files before relying on them, and apply their organization's safety, security, and compliance requirements before deployment.\n\nFile v1.0.14:skill.manifest.json\n\n{\n  \"references\": []\n}\n\nArchive v1.0.13: 4 files, 3673 bytes\n\nFiles: skill-card.md (2014b), skill.manifest.json (23b), SKILL.md (4613b), _meta.json (136b)\n\nFile v1.0.13:SKILL.md\n\n---\nname: stable-audio-2.5\ndescription: Use when someone wants light instrumental background music — an ambient bed under dialogue or underscore for reels and explainers.\nlicense: MIT\nmetadata:\n  version: \"1.0.13\"\n  package: pruna-skills\n  provider: replicate\n  replicate_model: stability-ai/stable-audio-2.5\n---\n\n## Prerequisites\n\nInstall and load these skills before generating (skip if already in context via `@pruna`):\n\n| Skill | Description | Install |\n| --- | --- | --- |\n| `generation-diversity` | Use when writing any generative prompt — ritual seed, explicit structure, scenario axes, and quality gates before paid API calls. | `npx skills add PrunaAI/pruna-skills@generation-diversity -y` |\n| `audio-prompting` | Use when crafting TTS, music, or bed prompts for any generative audio model — director style, song structure, and post-production layering. | `npx skills add PrunaAI/pruna-skills@audio-prompting -y` |\n| `pruna-api` | Use before any Pruna or Replicate HTTP call — credentials, upload/poll/download, parallel batches, and agent safety. | `npx skills add PrunaAI/pruna-skills@pruna-api -y` |\n\nOr install the full suite once: `npx skills add PrunaAI/pruna-skills@pruna -y`\n\nFollow each skill's **Before generating** / craft sections — do not restate guide content here.\n\n## Agent habit\n\nIn the **first reply**, name `` `stable-audio-2.5` `` in backticks, confirm `REPLICATE_API_TOKEN` (or stop with signup links from `pruna-api`), then ask for required inputs. Open intake → **`generation-diversity`** clarification intake before the first `POST`. Redirect when **When NOT to use** fits better.\n\n## When NOT to use\n\nUse a different skill instead:\n\n| Skill | Description | Install |\n| --- | --- | --- |\n| `music-2.5` | Use when someone wants an original AI song with vocals — sung lyrics, a style prompt track, or source audio for a music video. | `npx skills add PrunaAI/pruna-skills@music-2.5 -y` |\n| `gemini-3.1-flash-tts` | Use when someone needs spoken narration or voiceover — explainer tracks, documentary lines, or voice to pair with generated video. | `npx skills add PrunaAI/pruna-skills@gemini-3.1-flash-tts -y` |\n\n## Environment\n\n```bash\nexport REPLICATE_API_TOKEN=r8_...\n```\n\nRequires **`ffmpeg`** and **`ffprobe`** on PATH for the mix step.\n\n## HTTP (curl)\n\n```bash\ncurl -s -X POST \\\n  -H \"Authorization: Bearer ${REPLICATE_API_TOKEN}\" \\\n  -H \"Content-Type: application/json\" \\\n  -d '{\n    \"input\": {\n      \"prompt\": \"Instrumental light electronic pop bed, soft groove and mellow synth pads, calm positive tech atmosphere, understated background music, no vocals, 94 BPM\",\n      \"duration\": 90,\n      \"steps\": 8,\n      \"cfg_scale\": 1\n    }\n  }' \\\n  \"https://api.replicate.com/v1/models/stability-ai/stable-audio-2.5/predictions\"\n```\n\nPoll `urls.get` until `status` is `succeeded`; download `output` MP3.\n\n## Before generating\n\n1. Complete Prerequisites guide reading order — bed prompt craft: `audio-prompting` **Worked examples** (instrumental bed).\n2. Confirm **`prompt`**, **`duration`** (match or slightly exceed reel length), and mix **`volume`** (~0.08–0.15 under VO). When listing fields, name **`REPLICATE_API_TOKEN`** (Replicate — not `PRUNA_API_KEY`).\n3. **Model notes:** lead with **Instrumental** and **no vocals**. Duration 1–190s. Prefer understated beds (BPM ~88–98 for tech launch reels) so music does not compete with dialogue.\n\n## Required input\n\n- `prompt` (string)\n\n## Common optional fields\n\n- `duration` — seconds, 1–190\n- `steps` — 4–8 (default 8)\n- `cfg_scale` — 1–25 (default 1)\n- `seed` — optional integer\n\n## Typical next steps\n\nCommon follow-ons after this skill:\n\n| Skill | Description | Install |\n| --- | --- | --- |\n| `gemini-3.1-flash-tts` | Use when someone needs spoken narration or voiceover — explainer tracks, documentary lines, or voice to pair with generated video. | `npx skills add PrunaAI/pruna-skills@gemini-3.1-flash-tts -y` |\n| `visual-transition-reel` | Use when someone wants a montage with transitions between shots — action-sequence reel or multi-scene piece where narration is optional. | `npx skills add PrunaAI/pruna-skills@visual-transition-reel -y` |\n| `video-editing` | Use when assembling or polishing already-rendered clips with ffmpeg — concat, crossfades, burned captions and subtitles, text/logo overlays, before/after sliders, background music beds, platform export — or when composing a multi-layer HTML combination video with Hyperframes. Not for AI video generation, prompt craft, or model-based video edits. | `npx skills add PrunaAI/pruna-skills@video-editing -y` |\n\nFile v1.0.13:_meta.json\n\n{\n  \"ownerId\": \"kn7cagwf7q3t0cxrgteb7xk0bh81j0eb\",\n  \"slug\": \"stable-audio-2-5\",\n  \"version\": \"1.0.13\",\n  \"publishedAt\": 1789653549590\n}\n\nFile v1.0.13:skill-card.md\n\n## Description:\n\nUse when someone wants light instrumental background music - an ambient bed under dialogue or underscore for reels and explainers.\n\nThis skill is ready for commercial/non-commercial use.\n\n## Publisher:\n\n[pruna-ai](https://clawhub.ai/user/pruna-ai)\n\n### License/Terms of Use:\n\nMIT-0\n\n## Use Case:\n\nExternal creators and developers use this skill to guide an agent through generating light instrumental background music with Replicate's stability-ai/stable-audio-2.5 model, including prompt intake, API calling, polling, download, and basic mix guidance.\n\n### Deployment Geography for Use:\n\nGlobal\n\n## Known Risks and Mitigations:\n\nRisk: The security evidence says the skill asks users to install multiple unpinned follow-on skills that can change agent behavior later.\n\nMitigation: Review the referenced PrunaAI skill sources before installation and prefer pinned versions or reviewed commits where available.\n\nRisk: The skill sends generation prompts to Replicate and requires a Replicate API token.\n\nMitigation: Only provide REPLICATE_API_TOKEN in environments where sending prompts to Replicate is acceptable.\n\n## Reference(s):\n\n- [ClawHub skill page](https://clawhub.ai/pruna-ai/skills/stable-audio-2-5)\n- [Replicate Stable Audio 2.5 prediction API](https://api.replicate.com/v1/models/stability-ai/stable-audio-2.5/predictions)\n\n## Skill Output:\n\n**Output Type(s):** [text, markdown, shell commands, configuration, guidance]\n\n**Output Format:** [Markdown with inline shell commands and JSON/curl examples]\n\n**Output Parameters:** [1D]\n\n**Other Properties Related to Output:** [Requires a Replicate API token and ffmpeg/ffprobe for the mix step.]\n\n## Skill Version(s):\n\n1.0.13 (source: server release evidence and frontmatter)\n\n## Ethical Considerations:\n\nUsers should evaluate whether this skill is appropriate for their environment, review any generated or modified files before relying on them, and apply their organization's safety, security, and compliance requirements before deployment.\n\nFile v1.0.13:skill.manifest.json\n\n{\n  \"references\": []\n}\n\nArchive v1.0.12: 4 files, 3673 bytes\n\nFiles: skill-card.md (2026b), skill.manifest.json (23b), SKILL.md (4613b), _meta.json (136b)\n\nFile v1.0.12:SKILL.md\n\n---\nname: stable-audio-2.5\ndescription: Use when someone wants light instrumental background music — an ambient bed under dialogue or underscore for reels and explainers.\nlicense: MIT\nmetadata:\n  version: \"1.0.12\"\n  package: pruna-skills\n  provider: replicate\n  replicate_model: stability-ai/stable-audio-2.5\n---\n\n## Prerequisites\n\nInstall and load these skills before generating (skip if already in context via `@pruna`):\n\n| Skill | Description | Install |\n| --- | --- | --- |\n| `generation-diversity` | Use when writing any generative prompt — ritual seed, explicit structure, scenario axes, and quality gates before paid API calls. | `npx skills add PrunaAI/pruna-skills@generation-diversity -y` |\n| `audio-prompting` | Use when crafting TTS, music, or bed prompts for any generative audio model — director style, song structure, and post-production layering. | `npx skills add PrunaAI/pruna-skills@audio-prompting -y` |\n| `pruna-api` | Use before any Pruna or Replicate HTTP call — credentials, upload/poll/download, parallel batches, and agent safety. | `npx skills add PrunaAI/pruna-skills@pruna-api -y` |\n\nOr install the full suite once: `npx skills add PrunaAI/pruna-skills@pruna -y`\n\nFollow each skill's **Before generating** / craft sections — do not restate guide content here.\n\n## Agent habit\n\nIn the **first reply**, name `` `stable-audio-2.5` `` in backticks, confirm `REPLICATE_API_TOKEN` (or stop with signup links from `pruna-api`), then ask for required inputs. Open intake → **`generation-diversity`** clarification intake before the first `POST`. Redirect when **When NOT to use** fits better.\n\n## When NOT to use\n\nUse a different skill instead:\n\n| Skill | Description | Install |\n| --- | --- | --- |\n| `music-2.5` | Use when someone wants an original AI song with vocals — sung lyrics, a style prompt track, or source audio for a music video. | `npx skills add PrunaAI/pruna-skills@music-2.5 -y` |\n| `gemini-3.1-flash-tts` | Use when someone needs spoken narration or voiceover — explainer tracks, documentary lines, or voice to pair with generated video. | `npx skills add PrunaAI/pruna-skills@gemini-3.1-flash-tts -y` |\n\n## Environment\n\n```bash\nexport REPLICATE_API_TOKEN=r8_...\n```\n\nRequires **`ffmpeg`** and **`ffprobe`** on PATH for the mix step.\n\n## HTTP (curl)\n\n```bash\ncurl -s -X POST \\\n  -H \"Authorization: Bearer ${REPLICATE_API_TOKEN}\" \\\n  -H \"Content-Type: application/json\" \\\n  -d '{\n    \"input\": {\n      \"prompt\": \"Instrumental light electronic pop bed, soft groove and mellow synth pads, calm positive tech atmosphere, understated background music, no vocals, 94 BPM\",\n      \"duration\": 90,\n      \"steps\": 8,\n      \"cfg_scale\": 1\n    }\n  }' \\\n  \"https://api.replicate.com/v1/models/stability-ai/stable-audio-2.5/predictions\"\n```\n\nPoll `urls.get` until `status` is `succeeded`; download `output` MP3.\n\n## Before generating\n\n1. Complete Prerequisites guide reading order — bed prompt craft: `audio-prompting` **Worked examples** (instrumental bed).\n2. Confirm **`prompt`**, **`duration`** (match or slightly exceed reel length), and mix **`volume`** (~0.08–0.15 under VO). When listing fields, name **`REPLICATE_API_TOKEN`** (Replicate — not `PRUNA_API_KEY`).\n3. **Model notes:** lead with **Instrumental** and **no vocals**. Duration 1–190s. Prefer understated beds (BPM ~88–98 for tech launch reels) so music does not compete with dialogue.\n\n## Required input\n\n- `prompt` (string)\n\n## Common optional fields\n\n- `duration` — seconds, 1–190\n- `steps` — 4–8 (default 8)\n- `cfg_scale` — 1–25 (default 1)\n- `seed` — optional integer\n\n## Typical next steps\n\nCommon follow-ons after this skill:\n\n| Skill | Description | Install |\n| --- | --- | --- |\n| `gemini-3.1-flash-tts` | Use when someone needs spoken narration or voiceover — explainer tracks, documentary lines, or voice to pair with generated video. | `npx skills add PrunaAI/pruna-skills@gemini-3.1-flash-tts -y` |\n| `visual-transition-reel` | Use when someone wants a montage with transitions between shots — action-sequence reel or multi-scene piece where narration is optional. | `npx skills add PrunaAI/pruna-skills@visual-transition-reel -y` |\n| `video-editing` | Use when assembling or polishing already-rendered clips with ffmpeg — concat, crossfades, burned captions and subtitles, text/logo overlays, before/after sliders, background music beds, platform export — or when composing a multi-layer HTML combination video with Hyperframes. Not for AI video generation, prompt craft, or model-based video edits. | `npx skills add PrunaAI/pruna-skills@video-editing -y` |\n\nFile v1.0.12:_meta.json\n\n{\n  \"ownerId\": \"kn7cagwf7q3t0cxrgteb7xk0bh81j0eb\",\n  \"slug\": \"stable-audio-2-5\",\n  \"version\": \"1.0.12\",\n  \"publishedAt\": 1789048662529\n}\n\nFile v1.0.12:skill-card.md\n\n## Description:\n\nUse when someone wants light instrumental background music - an ambient bed under dialogue or underscore for reels and explainers.\n\nThis skill is ready for commercial/non-commercial use.\n\n## Publisher:\n\n[pruna-ai](https://clawhub.ai/user/pruna-ai)\n\n### License/Terms of Use:\n\nMIT-0\n\n## Use Case:\n\nDevelopers and content creators use this skill to prepare prompts, API calls, and mix guidance for generating light instrumental background music with Replicate's stability-ai/stable-audio-2.5 model.\n\n### Deployment Geography for Use:\n\nGlobal\n\n## Known Risks and Mitigations:\n\nRisk: Unpinned remote skill-install commands can change the agent environment with content that has not been reviewed for this release.\n\nMitigation: Install only required prerequisite skills, review commands before running them, and prefer pinned or verified installer and PrunaAI skill revisions.\n\nRisk: Prompts and generation inputs are sent to Replicate.\n\nMitigation: Avoid including sensitive data in prompts and use a scoped Replicate token rather than a broadly privileged token.\n\n## Reference(s):\n\n- [ClawHub skill page](https://clawhub.ai/pruna-ai/skills/stable-audio-2-5)\n- [Replicate Stable Audio 2.5 predictions API](https://api.replicate.com/v1/models/stability-ai/stable-audio-2.5/predictions)\n\n## Skill Output:\n\n**Output Type(s):** [guidance, shell commands, configuration, API calls]\n\n**Output Format:** [Markdown with inline bash and curl examples]\n\n**Output Parameters:** [1D]\n\n**Other Properties Related to Output:** [Requires REPLICATE_API_TOKEN plus ffmpeg and ffprobe for the mix step; generated audio is downloaded after polling the Replicate prediction output.]\n\n## Skill Version(s):\n\n1.0.12 (source: server release metadata and skill metadata)\n\n## Ethical Considerations:\n\nUsers should evaluate whether this skill is appropriate for their environment, review any generated or modified files before relying on them, and apply their organization's safety, security, and compliance requirements before deployment.\n\nFile v1.0.12:skill.manifest.json\n\n{\n  \"references\": []\n}\n\nArchive v1.0.11: 4 files, 3727 bytes\n\nFiles: skill-card.md (2109b), skill.manifest.json (23b), SKILL.md (4613b), _meta.json (136b)\n\nFile v1.0.11:SKILL.md\n\n---\nname: stable-audio-2.5\ndescription: Use when someone wants light instrumental background music — an ambient bed under dialogue or underscore for reels and explainers.\nlicense: MIT\nmetadata:\n  version: \"1.0.11\"\n  package: pruna-skills\n  provider: replicate\n  replicate_model: stability-ai/stable-audio-2.5\n---\n\n## Prerequisites\n\nInstall and load these skills before generating (skip if already in context via `@pruna`):\n\n| Skill | Description | Install |\n| --- | --- | --- |\n| `generation-diversity` | Use when writing any generative prompt — ritual seed, explicit structure, scenario axes, and quality gates before paid API calls. | `npx skills add PrunaAI/pruna-skills@generation-diversity -y` |\n| `audio-prompting` | Use when crafting TTS, music, or bed prompts for any generative audio model — director style, song structure, and post-production layering. | `npx skills add PrunaAI/pruna-skills@audio-prompting -y` |\n| `pruna-api` | Use before any Pruna or Replicate HTTP call — credentials, upload/poll/download, parallel batches, and agent safety. | `npx skills add PrunaAI/pruna-skills@pruna-api -y` |\n\nOr install the full suite once: `npx skills add PrunaAI/pruna-skills@pruna -y`\n\nFollow each skill's **Before generating** / craft sections — do not restate guide content here.\n\n## Agent habit\n\nIn the **first reply**, name `` `stable-audio-2.5` `` in backticks, confirm `REPLICATE_API_TOKEN` (or stop with signup links from `pruna-api`), then ask for required inputs. Open intake → **`generation-diversity`** clarification intake before the first `POST`. Redirect when **When NOT to use** fits better.\n\n## When NOT to use\n\nUse a different skill instead:\n\n| Skill | Description | Install |\n| --- | --- | --- |\n| `music-2.5` | Use when someone wants an original AI song with vocals — sung lyrics, a style prompt track, or source audio for a music video. | `npx skills add PrunaAI/pruna-skills@music-2.5 -y` |\n| `gemini-3.1-flash-tts` | Use when someone needs spoken narration or voiceover — explainer tracks, documentary lines, or voice to pair with generated video. | `npx skills add PrunaAI/pruna-skills@gemini-3.1-flash-tts -y` |\n\n## Environment\n\n```bash\nexport REPLICATE_API_TOKEN=r8_...\n```\n\nRequires **`ffmpeg`** and **`ffprobe`** on PATH for the mix step.\n\n## HTTP (curl)\n\n```bash\ncurl -s -X POST \\\n  -H \"Authorization: Bearer ${REPLICATE_API_TOKEN}\" \\\n  -H \"Content-Type: application/json\" \\\n  -d '{\n    \"input\": {\n      \"prompt\": \"Instrumental light electronic pop bed, soft groove and mellow synth pads, calm positive tech atmosphere, understated background music, no vocals, 94 BPM\",\n      \"duration\": 90,\n      \"steps\": 8,\n      \"cfg_scale\": 1\n    }\n  }' \\\n  \"https://api.replicate.com/v1/models/stability-ai/stable-audio-2.5/predictions\"\n```\n\nPoll `urls.get` until `status` is `succeeded`; download `output` MP3.\n\n## Before generating\n\n1. Complete Prerequisites guide reading order — bed prompt craft: `audio-prompting` **Worked examples** (instrumental bed).\n2. Confirm **`prompt`**, **`duration`** (match or slightly exceed reel length), and mix **`volume`** (~0.08–0.15 under VO). When listing fields, name **`REPLICATE_API_TOKEN`** (Replicate — not `PRUNA_API_KEY`).\n3. **Model notes:** lead with **Instrumental** and **no vocals**. Duration 1–190s. Prefer understated beds (BPM ~88–98 for tech launch reels) so music does not compete with dialogue.\n\n## Required input\n\n- `prompt` (string)\n\n## Common optional fields\n\n- `duration` — seconds, 1–190\n- `steps` — 4–8 (default 8)\n- `cfg_scale` — 1–25 (default 1)\n- `seed` — optional integer\n\n## Typical next steps\n\nCommon follow-ons after this skill:\n\n| Skill | Description | Install |\n| --- | --- | --- |\n| `gemini-3.1-flash-tts` | Use when someone needs spoken narration or voiceover — explainer tracks, documentary lines, or voice to pair with generated video. | `npx skills add PrunaAI/pruna-skills@gemini-3.1-flash-tts -y` |\n| `visual-transition-reel` | Use when someone wants a montage with transitions between shots — action-sequence reel or multi-scene piece where narration is optional. | `npx skills add PrunaAI/pruna-skills@visual-transition-reel -y` |\n| `video-editing` | Use when assembling or polishing already-rendered clips with ffmpeg — concat, crossfades, burned captions and subtitles, text/logo overlays, before/after sliders, background music beds, platform export — or when composing a multi-layer HTML combination video with Hyperframes. Not for AI video generation, prompt craft, or model-based video edits. | `npx skills add PrunaAI/pruna-skills@video-editing -y` |\n\nFile v1.0.11:_meta.json\n\n{\n  \"ownerId\": \"kn7cagwf7q3t0cxrgteb7xk0bh81j0eb\",\n  \"slug\": \"stable-audio-2-5\",\n  \"version\": \"1.0.11\",\n  \"publishedAt\": 1788444717797\n}\n\nFile v1.0.11:skill-card.md\n\n## Description:\n\nUse when someone wants light instrumental background music, such as an ambient bed under dialogue or underscore for reels and explainers.\n\nThis skill is ready for commercial/non-commercial use.\n\n## Publisher:\n\n[pruna-ai](https://clawhub.ai/user/pruna-ai)\n\n### License/Terms of Use:\n\nMIT-0\n\n## Use Case:\n\nCreators, marketers, and developers use this skill to craft Replicate Stable Audio 2.5 prompts and API calls for light instrumental background beds, then download and mix generated MP3 output under narration or video.\n\n### Deployment Geography for Use:\n\nGlobal\n\n## Known Risks and Mitigations:\n\nRisk: Remote prerequisite install commands can add or update skills from a remote repository.\n\nMitigation: Review each npx skills add command and installed skill content before running generation workflows.\n\nRisk: Secrets or sensitive private content may be exposed if included in prompts sent to Replicate.\n\nMitigation: Do not place secrets or sensitive private content in audio prompts.\n\nRisk: Missing local media tools can block the mix step.\n\nMitigation: Confirm ffmpeg and ffprobe are installed on PATH before planning a mixed audio deliverable.\n\n## Reference(s):\n\n- [ClawHub skill page](https://clawhub.ai/pruna-ai/skills/stable-audio-2-5)\n- [Replicate Stable Audio 2.5 prediction API](https://api.replicate.com/v1/models/stability-ai/stable-audio-2.5/predictions)\n\n## Skill Output:\n\n**Output Type(s):** [text, markdown, shell commands, configuration, guidance]\n\n**Output Format:** [Markdown guidance with shell commands and JSON API request examples]\n\n**Output Parameters:** [1D]\n\n**Other Properties Related to Output:** [Guides use of REPLICATE_API_TOKEN, Replicate prediction polling, MP3 download, and optional ffmpeg-based mixing.]\n\n## Skill Version(s):\n\n1.0.11 (source: server release evidence and skill metadata)\n\n## Ethical Considerations:\n\nUsers should evaluate whether this skill is appropriate for their environment, review any generated or modified files before relying on them, and apply their organization's safety, security, and compliance requirements before deployment.\n\nFile v1.0.11:skill.manifest.json\n\n{\n  \"references\": []\n}\n\nArchive v1.0.10: 4 files, 3702 bytes\n\nFiles: skill-card.md (2129b), skill.manifest.json (23b), SKILL.md (4613b), _meta.json (136b)\n\nFile v1.0.10:SKILL.md\n\n---\nname: stable-audio-2.5\ndescription: Use when someone wants light instrumental background music — an ambient bed under dialogue or underscore for reels and explainers.\nlicense: MIT\nmetadata:\n  version: \"1.0.10\"\n  package: pruna-skills\n  provider: replicate\n  replicate_model: stability-ai/stable-audio-2.5\n---\n\n## Prerequisites\n\nInstall and load these skills before generating (skip if already in context via `@pruna`):\n\n| Skill | Description | Install |\n| --- | --- | --- |\n| `generation-diversity` | Use when writing any generative prompt — ritual seed, explicit structure, scenario axes, and quality gates before paid API calls. | `npx skills add PrunaAI/pruna-skills@generation-diversity -y` |\n| `audio-prompting` | Use when crafting TTS, music, or bed prompts for any generative audio model — director style, song structure, and post-production layering. | `npx skills add PrunaAI/pruna-skills@audio-prompting -y` |\n| `pruna-api` | Use before any Pruna or Replicate HTTP call — credentials, upload/poll/download, parallel batches, and agent safety. | `npx skills add PrunaAI/pruna-skills@pruna-api -y` |\n\nOr install the full suite once: `npx skills add PrunaAI/pruna-skills@pruna -y`\n\nFollow each skill's **Before generating** / craft sections — do not restate guide content here.\n\n## Agent habit\n\nIn the **first reply**, name `` `stable-audio-2.5` `` in backticks, confirm `REPLICATE_API_TOKEN` (or stop with signup links from `pruna-api`), then ask for required inputs. Open intake → **`generation-diversity`** clarification intake before the first `POST`. Redirect when **When NOT to use** fits better.\n\n## When NOT to use\n\nUse a different skill instead:\n\n| Skill | Description | Install |\n| --- | --- | --- |\n| `music-2.5` | Use when someone wants an original AI song with vocals — sung lyrics, a style prompt track, or source audio for a music video. | `npx skills add PrunaAI/pruna-skills@music-2.5 -y` |\n| `gemini-3.1-flash-tts` | Use when someone needs spoken narration or voiceover — explainer tracks, documentary lines, or voice to pair with generated video. | `npx skills add PrunaAI/pruna-skills@gemini-3.1-flash-tts -y` |\n\n## Environment\n\n```bash\nexport REPLICATE_API_TOKEN=r8_...\n```\n\nRequires **`ffmpeg`** and **`ffprobe`** on PATH for the mix step.\n\n## HTTP (curl)\n\n```bash\ncurl -s -X POST \\\n  -H \"Authorization: Bearer ${REPLICATE_API_TOKEN}\" \\\n  -H \"Content-Type: application/json\" \\\n  -d '{\n    \"input\": {\n      \"prompt\": \"Instrumental light electronic pop bed, soft groove and mellow synth pads, calm positive tech atmosphere, understated background music, no vocals, 94 BPM\",\n      \"duration\": 90,\n      \"steps\": 8,\n      \"cfg_scale\": 1\n    }\n  }' \\\n  \"https://api.replicate.com/v1/models/stability-ai/stable-audio-2.5/predictions\"\n```\n\nPoll `urls.get` until `status` is `succeeded`; download `output` MP3.\n\n## Before generating\n\n1. Complete Prerequisites guide reading order — bed prompt craft: `audio-prompting` **Worked examples** (instrumental bed).\n2. Confirm **`prompt`**, **`duration`** (match or slightly exceed reel length), and mix **`volume`** (~0.08–0.15 under VO). When listing fields, name **`REPLICATE_API_TOKEN`** (Replicate — not `PRUNA_API_KEY`).\n3. **Model notes:** lead with **Instrumental** and **no vocals**. Duration 1–190s. Prefer understated beds (BPM ~88–98 for tech launch reels) so music does not compete with dialogue.\n\n## Required input\n\n- `prompt` (string)\n\n## Common optional fields\n\n- `duration` — seconds, 1–190\n- `steps` — 4–8 (default 8)\n- `cfg_scale` — 1–25 (default 1)\n- `seed` — optional integer\n\n## Typical next steps\n\nCommon follow-ons after this skill:\n\n| Skill | Description | Install |\n| --- | --- | --- |\n| `gemini-3.1-flash-tts` | Use when someone needs spoken narration or voiceover — explainer tracks, documentary lines, or voice to pair with generated video. | `npx skills add PrunaAI/pruna-skills@gemini-3.1-flash-tts -y` |\n| `visual-transition-reel` | Use when someone wants a montage with transitions between shots — action-sequence reel or multi-scene piece where narration is optional. | `npx skills add PrunaAI/pruna-skills@visual-transition-reel -y` |\n| `video-editing` | Use when assembling or polishing already-rendered clips with ffmpeg — concat, crossfades, burned captions and subtitles, text/logo overlays, before/after sliders, background music beds, platform export — or when composing a multi-layer HTML combination video with Hyperframes. Not for AI video generation, prompt craft, or model-based video edits. | `npx skills add PrunaAI/pruna-skills@video-editing -y` |\n\nFile v1.0.10:_meta.json\n\n{\n  \"ownerId\": \"kn7cagwf7q3t0cxrgteb7xk0bh81j0eb\",\n  \"slug\": \"stable-audio-2-5\",\n  \"version\": \"1.0.10\",\n  \"publishedAt\": 1787903867430\n}\n\nFile v1.0.10:skill-card.md\n\n## Description:\n\nUse when someone wants light instrumental background music - an ambient bed under dialogue or underscore for reels and explainers.\n\nThis skill is ready for commercial/non-commercial use.\n\n## Publisher:\n\n[pruna-ai](https://clawhub.ai/user/pruna-ai)\n\n### License/Terms of Use:\n\nMIT-0\n\n## Use Case:\n\nDevelopers and content creators use this skill to guide an agent through generating light instrumental background music with the Replicate-hosted stability-ai/stable-audio-2.5 model.\n\n### Deployment Geography for Use:\n\nGlobal\n\n## Known Risks and Mitigations:\n\nRisk: Prompts and generation settings are sent to Replicate, and use may consume paid API credits.\n\nMitigation: Confirm the Replicate API token, review prompt content before requests, and verify expected cost or credit usage before generation.\n\nRisk: The skill depends on local ffmpeg and ffprobe availability for the mix step.\n\nMitigation: Verify ffmpeg and ffprobe are installed and on PATH before using the generated audio in a mix workflow.\n\nRisk: Recommended prerequisite skills are not included in this artifact.\n\nMitigation: Review and install the referenced Pruna prerequisite skills separately before following their guidance.\n\n## Reference(s):\n\n- [ClawHub skill page](https://clawhub.ai/pruna-ai/skills/stable-audio-2-5)\n- [Replicate Stable Audio 2.5 predictions endpoint](https://api.replicate.com/v1/models/stability-ai/stable-audio-2.5/predictions)\n\n## Skill Output:\n\n**Output Type(s):** [Guidance, Markdown, Shell commands, Configuration, API calls]\n\n**Output Format:** [Markdown with inline bash and curl examples]\n\n**Output Parameters:** [1D]\n\n**Other Properties Related to Output:** [Guides prompt, duration, generation settings, polling, MP3 download, and ffmpeg-based mix preparation.]\n\n## Skill Version(s):\n\n1.0.10 (source: server release metadata and skill metadata)\n\n## Ethical Considerations:\n\nUsers should evaluate whether this skill is appropriate for their environment, review any generated or modified files before relying on them, and apply their organization's safety, security, and compliance requirements before deployment.\n\nFile v1.0.10:skill.manifest.json\n\n{\n  \"references\": []\n}\n\nArchive v1.0.9: 4 files, 3772 bytes\n\nFiles: skill-card.md (2299b), skill.manifest.json (23b), SKILL.md (4612b), _meta.json (135b)\n\nFile v1.0.9:SKILL.md\n\n---\nname: stable-audio-2.5\ndescription: Use when someone wants light instrumental background music — an ambient bed under dialogue or underscore for reels and explainers.\nlicense: MIT\nmetadata:\n  version: \"1.0.9\"\n  package: pruna-skills\n  provider: replicate\n  replicate_model: stability-ai/stable-audio-2.5\n---\n\n## Prerequisites\n\nInstall and load these skills before generating (skip if already in context via `@pruna`):\n\n| Skill | Description | Install |\n| --- | --- | --- |\n| `generation-diversity` | Use when writing any generative prompt — ritual seed, explicit structure, scenario axes, and quality gates before paid API calls. | `npx skills add PrunaAI/pruna-skills@generation-diversity -y` |\n| `audio-prompting` | Use when crafting TTS, music, or bed prompts for any generative audio model — director style, song structure, and post-production layering. | `npx skills add PrunaAI/pruna-skills@audio-prompting -y` |\n| `pruna-api` | Use before any Pruna or Replicate HTTP call — credentials, upload/poll/download, parallel batches, and agent safety. | `npx skills add PrunaAI/pruna-skills@pruna-api -y` |\n\nOr install the full suite once: `npx skills add PrunaAI/pruna-skills@pruna -y`\n\nFollow each skill's **Before generating** / craft sections — do not restate guide content here.\n\n## Agent habit\n\nIn the **first reply**, name `` `stable-audio-2.5` `` in backticks, confirm `REPLICATE_API_TOKEN` (or stop with signup links from `pruna-api`), then ask for required inputs. Open intake → **`generation-diversity`** clarification intake before the first `POST`. Redirect when **When NOT to use** fits better.\n\n## When NOT to use\n\nUse a different skill instead:\n\n| Skill | Description | Install |\n| --- | --- | --- |\n| `music-2.5` | Use when someone wants an original AI song with vocals — sung lyrics, a style prompt track, or source audio for a music video. | `npx skills add PrunaAI/pruna-skills@music-2.5 -y` |\n| `gemini-3.1-flash-tts` | Use when someone needs spoken narration or voiceover — explainer tracks, documentary lines, or voice to pair with generated video. | `npx skills add PrunaAI/pruna-skills@gemini-3.1-flash-tts -y` |\n\n## Environment\n\n```bash\nexport REPLICATE_API_TOKEN=r8_...\n```\n\nRequires **`ffmpeg`** and **`ffprobe`** on PATH for the mix step.\n\n## HTTP (curl)\n\n```bash\ncurl -s -X POST \\\n  -H \"Authorization: Bearer ${REPLICATE_API_TOKEN}\" \\\n  -H \"Content-Type: application/json\" \\\n  -d '{\n    \"input\": {\n      \"prompt\": \"Instrumental light electronic pop bed, soft groove and mellow synth pads, calm positive tech atmosphere, understated background music, no vocals, 94 BPM\",\n      \"duration\": 90,\n      \"steps\": 8,\n      \"cfg_scale\": 1\n    }\n  }' \\\n  \"https://api.replicate.com/v1/models/stability-ai/stable-audio-2.5/predictions\"\n```\n\nPoll `urls.get` until `status` is `succeeded`; download `output` MP3.\n\n## Before generating\n\n1. Complete Prerequisites guide reading order — bed prompt craft: `audio-prompting` **Worked examples** (instrumental bed).\n2. Confirm **`prompt`**, **`duration`** (match or slightly exceed reel length), and mix **`volume`** (~0.08–0.15 under VO). When listing fields, name **`REPLICATE_API_TOKEN`** (Replicate — not `PRUNA_API_KEY`).\n3. **Model notes:** lead with **Instrumental** and **no vocals**. Duration 1–190s. Prefer understated beds (BPM ~88–98 for tech launch reels) so music does not compete with dialogue.\n\n## Required input\n\n- `prompt` (string)\n\n## Common optional fields\n\n- `duration` — seconds, 1–190\n- `steps` — 4–8 (default 8)\n- `cfg_scale` — 1–25 (default 1)\n- `seed` — optional integer\n\n## Typical next steps\n\nCommon follow-ons after this skill:\n\n| Skill | Description | Install |\n| --- | --- | --- |\n| `gemini-3.1-flash-tts` | Use when someone needs spoken narration or voiceover — explainer tracks, documentary lines, or voice to pair with generated video. | `npx skills add PrunaAI/pruna-skills@gemini-3.1-flash-tts -y` |\n| `visual-transition-reel` | Use when someone wants a montage with transitions between shots — action-sequence reel or multi-scene piece where narration is optional. | `npx skills add PrunaAI/pruna-skills@visual-transition-reel -y` |\n| `video-editing` | Use when assembling or polishing already-rendered clips with ffmpeg — concat, crossfades, burned captions and subtitles, text/logo overlays, before/after sliders, background music beds, platform export — or when composing a multi-layer HTML combination video with Hyperframes. Not for AI video generation, prompt craft, or model-based video edits. | `npx skills add PrunaAI/pruna-skills@video-editing -y` |\n\nFile v1.0.9:_meta.json\n\n{\n  \"ownerId\": \"kn7cagwf7q3t0cxrgteb7xk0bh81j0eb\",\n  \"slug\": \"stable-audio-2-5\",\n  \"version\": \"1.0.9\",\n  \"publishedAt\": 1785824309963\n}\n\nFile v1.0.9:skill-card.md\n\n## Description: <br>\nUse when someone wants light instrumental background music -- an ambient bed under dialogue or underscore for reels and explainers. <br>\n\nThis skill is ready for commercial/non-commercial use. <br>\n\n## Publisher: <br>\n[pruna-ai](https://clawhub.ai/user/pruna-ai) <br>\n\n### License/Terms of Use: <br>\nMIT-0 <br>\n\n\n## Use Case: <br>\nDevelopers, creators, and agent operators use this skill to generate light instrumental background music through Replicate's Stable Audio 2.5 model for reels, explainers, and dialogue beds. <br>\n\n### Deployment Geography for Use: <br>\nGlobal <br>\n\n## Known Risks and Mitigations: <br>\nRisk: Using the skill requires a Replicate API token and may incur Replicate usage costs. <br>\nMitigation: Confirm REPLICATE_API_TOKEN is available, keep it out of prompts and logs, and run generation only after the user accepts possible provider costs. <br>\nRisk: The skill delegates prompt-crafting and API-handling guidance to prerequisite Pruna skills. <br>\nMitigation: Review and load the referenced prerequisite skills before making paid API calls or generating audio. <br>\nRisk: The audio mix step depends on ffmpeg and ffprobe being installed on PATH. <br>\nMitigation: Check tool availability before attempting the mix step and stop with setup guidance if either dependency is missing. <br>\n\n\n## Reference(s): <br>\n- [ClawHub skill page](https://clawhub.ai/pruna-ai/skills/stable-audio-2-5) <br>\n- [Replicate Stable Audio 2.5 prediction endpoint](https://api.replicate.com/v1/models/stability-ai/stable-audio-2.5/predictions) <br>\n\n\n## Skill Output: <br>\n**Output Type(s):** [Guidance, Shell commands, Configuration, API calls] <br>\n**Output Format:** [Markdown with inline bash and curl examples] <br>\n**Output Parameters:** [1D] <br>\n**Other Properties Related to Output:** [Guides a Replicate prediction request, polling, MP3 download, and optional audio mix preparation.] <br>\n\n## Skill Version(s): <br>\n1.0.9 (source: server release metadata and SKILL.md frontmatter) <br>\n\n## Ethical Considerations: <br>\nUsers should evaluate whether this skill is appropriate for their environment, review any generated or modified files before relying on them, and apply their organization's safety, security, and compliance requirements before deployment. <br>\n\nFile v1.0.9:skill.manifest.json\n\n{\n  \"references\": []\n}\n\nArchive v1.0.8: 4 files, 3703 bytes\n\nFiles: skill-card.md (2223b), skill.manifest.json (23b), SKILL.md (4612b), _meta.json (135b)\n\nFile v1.0.8:SKILL.md\n\n---\nname: stable-audio-2.5\ndescription: Use when someone wants light instrumental background music — an ambient bed under dialogue or underscore for reels and explainers.\nlicense: MIT\nmetadata:\n  version: \"1.0.8\"\n  package: pruna-skills\n  provider: replicate\n  replicate_model: stability-ai/stable-audio-2.5\n---\n\n## Prerequisites\n\nInstall and load these skills before generating (skip if already in context via `@pruna`):\n\n| Skill | Description | Install |\n| --- | --- | --- |\n| `generation-diversity` | Use when writing any generative prompt — ritual seed, explicit structure, scenario axes, and quality gates before paid API calls. | `npx skills add PrunaAI/pruna-skills@generation-diversity -y` |\n| `audio-prompting` | Use when crafting TTS, music, or bed prompts for any generative audio model — director style, song structure, and post-production layering. | `npx skills add PrunaAI/pruna-skills@audio-prompting -y` |\n| `pruna-api` | Use before any Pruna or Replicate HTTP call — credentials, upload/poll/download, parallel batches, and agent safety. | `npx skills add PrunaAI/pruna-skills@pruna-api -y` |\n\nOr install the full suite once: `npx skills add PrunaAI/pruna-skills@pruna -y`\n\nFollow each skill's **Before generating** / craft sections — do not restate guide content here.\n\n## Agent habit\n\nIn the **first reply**, name `` `stable-audio-2.5` `` in backticks, confirm `REPLICATE_API_TOKEN` (or stop with signup links from `pruna-api`), then ask for required inputs. Open intake → **`generation-diversity`** clarification intake before the first `POST`. Redirect when **When NOT to use** fits better.\n\n## When NOT to use\n\nUse a different skill instead:\n\n| Skill | Description | Install |\n| --- | --- | --- |\n| `music-2.5` | Use when someone wants an original AI song with vocals — sung lyrics, a style prompt track, or source audio for a music video. | `npx skills add PrunaAI/pruna-skills@music-2.5 -y` |\n| `gemini-3.1-flash-tts` | Use when someone needs spoken narration or voiceover — explainer tracks, documentary lines, or voice to pair with generated video. | `npx skills add PrunaAI/pruna-skills@gemini-3.1-flash-tts -y` |\n\n## Environment\n\n```bash\nexport REPLICATE_API_TOKEN=r8_...\n```\n\nRequires **`ffmpeg`** and **`ffprobe`** on PATH for the mix step.\n\n## HTTP (curl)\n\n```bash\ncurl -s -X POST \\\n  -H \"Authorization: Bearer ${REPLICATE_API_TOKEN}\" \\\n  -H \"Content-Type: application/json\" \\\n  -d '{\n    \"input\": {\n      \"prompt\": \"Instrumental light electronic pop bed, soft groove and mellow synth pads, calm positive tech atmosphere, understated background music, no vocals, 94 BPM\",\n      \"duration\": 90,\n      \"steps\": 8,\n      \"cfg_scale\": 1\n    }\n  }' \\\n  \"https://api.replicate.com/v1/models/stability-ai/stable-audio-2.5/predictions\"\n```\n\nPoll `urls.get` until `status` is `succeeded`; download `output` MP3.\n\n## Before generating\n\n1. Complete Prerequisites guide reading order — bed prompt craft: `audio-prompting` **Worked examples** (instrumental bed).\n2. Confirm **`prompt`**, **`duration`** (match or slightly exceed reel length), and mix **`volume`** (~0.08–0.15 under VO). When listing fields, name **`REPLICATE_API_TOKEN`** (Replicate — not `PRUNA_API_KEY`).\n3. **Model notes:** lead with **Instrumental** and **no vocals**. Duration 1–190s. Prefer understated beds (BPM ~88–98 for tech launch reels) so music does not compete with dialogue.\n\n## Required input\n\n- `prompt` (string)\n\n## Common optional fields\n\n- `duration` — seconds, 1–190\n- `steps` — 4–8 (default 8)\n- `cfg_scale` — 1–25 (default 1)\n- `seed` — optional integer\n\n## Typical next steps\n\nCommon follow-ons after this skill:\n\n| Skill | Description | Install |\n| --- | --- | --- |\n| `gemini-3.1-flash-tts` | Use when someone needs spoken narration or voiceover — explainer tracks, documentary lines, or voice to pair with generated video. | `npx skills add PrunaAI/pruna-skills@gemini-3.1-flash-tts -y` |\n| `visual-transition-reel` | Use when someone wants a montage with transitions between shots — action-sequence reel or multi-scene piece where narration is optional. | `npx skills add PrunaAI/pruna-skills@visual-transition-reel -y` |\n| `video-editing` | Use when assembling or polishing already-rendered clips with ffmpeg — concat, crossfades, burned captions and subtitles, text/logo overlays, before/after sliders, background music beds, platform export — or when composing a multi-layer HTML combination video with Hyperframes. Not for AI video generation, prompt craft, or model-based video edits. | `npx skills add PrunaAI/pruna-skills@video-editing -y` |\n\nFile v1.0.8:_meta.json\n\n{\n  \"ownerId\": \"kn7cagwf7q3t0cxrgteb7xk0bh81j0eb\",\n  \"slug\": \"stable-audio-2-5\",\n  \"version\": \"1.0.8\",\n  \"publishedAt\": 1785259234327\n}\n\nFile v1.0.8:skill-card.md\n\n## Description: <br>\nHelps agents generate light instrumental background music with Stable Audio 2.5 on Replicate, such as ambient beds under dialogue or underscores for reels and explainers. <br>\n\nThis skill is ready for commercial/non-commercial use. <br>\n\n## Publisher: <br>\n[pruna-ai](https://clawhub.ai/user/pruna-ai) <br>\n\n### License/Terms of Use: <br>\nMIT-0 <br>\n\n\n## Use Case: <br>\nDevelopers and creators use this skill to prompt Stable Audio 2.5 for short instrumental background beds, confirm required inputs and credentials, submit a Replicate prediction, poll for completion, and download the generated MP3. <br>\n\n### Deployment Geography for Use: <br>\nGlobal <br>\n\n## Known Risks and Mitigations: <br>\nRisk: The skill depends on additional PrunaAI skills that can expand what is installed or loaded for the agent. <br>\nMitigation: Review the referenced dependency skills before installing or using this skill. <br>\nRisk: Prompts and generated audio are sent through Replicate when the API request is made. <br>\nMitigation: Use the skill only when that data flow is acceptable for the content being generated. <br>\nRisk: A Replicate API token is required for generation. <br>\nMitigation: Keep the token in the environment and do not hardcode it in files. <br>\n\n\n## Reference(s): <br>\n- [ClawHub skill page](https://clawhub.ai/pruna-ai/skills/stable-audio-2-5) <br>\n- [Replicate Stable Audio 2.5 prediction endpoint](https://api.replicate.com/v1/models/stability-ai/stable-audio-2.5/predictions) <br>\n\n\n## Skill Output: <br>\n**Output Type(s):** [text, markdown, shell commands, configuration, guidance] <br>\n**Output Format:** [Markdown guidance with bash and curl examples] <br>\n**Output Parameters:** [1D] <br>\n**Other Properties Related to Output:** [Requires a Replicate API token for generation and ffmpeg/ffprobe for the mix step.] <br>\n\n## Skill Version(s): <br>\n1.0.8 (source: server release evidence and skill metadata) <br>\n\n## Ethical Considerations: <br>\nUsers should evaluate whether this skill is appropriate for their environment, review any generated or modified files before relying on them, and apply their organization's safety, security, and compliance requirements before deployment. <br>\n\nFile v1.0.8:skill.manifest.json\n\n{\n  \"references\": []\n}\n\nArchive v1.0.7: 4 files, 3730 bytes\n\nFiles: skill-card.md (2317b), skill.manifest.json (23b), SKILL.md (4523b), _meta.json (135b)\n\nFile v1.0.7:SKILL.md\n\n---\nname: stable-audio-2.5\ndescription: Use when someone wants light instrumental background music — an ambient bed under dialogue or underscore for reels and explainers.\nlicense: MIT\nmetadata:\n  version: \"1.0.7\"\n  package: pruna-skills\n  provider: replicate\n  replicate_model: stability-ai/stable-audio-2.5\n---\n\n## Prerequisites\n\nInstall and load these skills before generating (skip if already in context via `@pruna`):\n\n| Skill | Description | Install |\n| --- | --- | --- |\n| `generation-diversity` | Use when writing any generative prompt — ritual seed, explicit structure, scenario axes, and quality gates before paid API calls. | `npx skills add PrunaAI/pruna-skills@generation-diversity -y` |\n| `audio-prompting` | Use when crafting TTS, music, or bed prompts for any generative audio model — director style, song structure, and post-production layering. | `npx skills add PrunaAI/pruna-skills@audio-prompting -y` |\n| `pruna-api` | Use before any Pruna or Replicate HTTP call — credentials, upload/poll/download, parallel batches, and agent safety. | `npx skills add PrunaAI/pruna-skills@pruna-api -y` |\n\nOr install the full suite once: `npx skills add PrunaAI/pruna-skills@pruna -y`\n\nFollow each skill's **Before generating** / craft sections — do not restate guide content here.\n\n## Agent habit\n\nIn the **first reply**, name `` `stable-audio-2.5` `` in backticks, confirm `REPLICATE_API_TOKEN` (or stop with signup links from `pruna-api`), then ask for required inputs. Redirect when **When NOT to use** fits better.\n\n## When NOT to use\n\nUse a different skill instead:\n\n| Skill | Description | Install |\n| --- | --- | --- |\n| `music-2.5` | Use when someone wants an original AI song with vocals — sung lyrics, a style prompt track, or source audio for a music video. | `npx skills add PrunaAI/pruna-skills@music-2.5 -y` |\n| `gemini-3.1-flash-tts` | Use when someone needs spoken narration or voiceover — explainer tracks, documentary lines, or voice to pair with generated video. | `npx skills add PrunaAI/pruna-skills@gemini-3.1-flash-tts -y` |\n\n## Environment\n\n```bash\nexport REPLICATE_API_TOKEN=r8_...\n```\n\nRequires **`ffmpeg`** and **`ffprobe`** on PATH for the mix step.\n\n## HTTP (curl)\n\n```bash\ncurl -s -X POST \\\n  -H \"Authorization: Bearer ${REPLICATE_API_TOKEN}\" \\\n  -H \"Content-Type: application/json\" \\\n  -d '{\n    \"input\": {\n      \"prompt\": \"Instrumental light electronic pop bed, soft groove and mellow synth pads, calm positive tech atmosphere, understated background music, no vocals, 94 BPM\",\n      \"duration\": 90,\n      \"steps\": 8,\n      \"cfg_scale\": 1\n    }\n  }' \\\n  \"https://api.replicate.com/v1/models/stability-ai/stable-audio-2.5/predictions\"\n```\n\nPoll `urls.get` until `status` is `succeeded`; download `output` MP3.\n\n## Before generating\n\n1. Complete Prerequisites guide reading order — bed prompt craft: `audio-prompting` **Worked examples** (instrumental bed).\n2. Confirm **`prompt`**, **`duration`** (match or slightly exceed reel length), and mix **`volume`** (~0.08–0.15 under VO). When listing fields, name **`REPLICATE_API_TOKEN`** (Replicate — not `PRUNA_API_KEY`).\n3. **Model notes:** lead with **Instrumental** and **no vocals**. Duration 1–190s. Prefer understated beds (BPM ~88–98 for tech launch reels) so music does not compete with dialogue.\n\n## Required input\n\n- `prompt` (string)\n\n## Common optional fields\n\n- `duration` — seconds, 1–190\n- `steps` — 4–8 (default 8)\n- `cfg_scale` — 1–25 (default 1)\n- `seed` — optional integer\n\n## Typical next steps\n\nCommon follow-ons after this skill:\n\n| Skill | Description | Install |\n| --- | --- | --- |\n| `gemini-3.1-flash-tts` | Use when someone needs spoken narration or voiceover — explainer tracks, documentary lines, or voice to pair with generated video. | `npx skills add PrunaAI/pruna-skills@gemini-3.1-flash-tts -y` |\n| `visual-transition-reel` | Use when someone wants a montage with transitions between shots — action-sequence reel or multi-scene piece where narration is optional. | `npx skills add PrunaAI/pruna-skills@visual-transition-reel -y` |\n| `video-editing` | Use when assembling or polishing already-rendered clips with ffmpeg — concat, crossfades, burned captions and subtitles, text/logo overlays, before/after sliders, background music beds, platform export — or when composing a multi-layer HTML combination video with Hyperframes. Not for AI video generation, prompt craft, or model-based video edits. | `npx skills add PrunaAI/pruna-skills@video-editing -y` |\n\nFile v1.0.7:_meta.json\n\n{\n  \"ownerId\": \"kn7cagwf7q3t0cxrgteb7xk0bh81j0eb\",\n  \"slug\": \"stable-audio-2-5\",\n  \"version\": \"1.0.7\",\n  \"publishedAt\": 1784810136093\n}\n\nFile v1.0.7:skill-card.md\n\n## Description: <br>\nUse when someone wants light instrumental background music, such as an ambient bed under dialogue or underscore for reels and explainers. <br>\n\nThis skill is ready for commercial/non-commercial use. <br>\n\n## Publisher: <br>\n[pruna-ai](https://clawhub.ai/user/pruna-ai) <br>\n\n### License/Terms of Use: <br>\nMIT-0 <br>\n\n\n## Use Case: <br>\nExternal developers and content-production agents use this skill to prepare prompts, environment setup, and Replicate API calls for generating understated instrumental background music. <br>\n\n### Deployment Geography for Use: <br>\nGlobal <br>\n\n## Known Risks and Mitigations: <br>\nRisk: Replicate API use may consume paid or quota-limited generation capacity. <br>\nMitigation: Confirm the prompt, duration, steps, and cfg_scale before making prediction requests. <br>\nRisk: API tokens can be exposed if pasted into shared prompts, logs, or generated text. <br>\nMitigation: Keep REPLICATE_API_TOKEN in the environment and avoid writing token values into shared text. <br>\nRisk: Generated music can interfere with dialogue or include unwanted vocal elements if the prompt is underspecified. <br>\nMitigation: Lead prompts with instrumental and no vocals, and keep mix volume around the documented background-bed range. <br>\n\n\n## Reference(s): <br>\n- [ClawHub skill page](https://clawhub.ai/pruna-ai/skills/stable-audio-2-5) <br>\n- [Replicate Stable Audio 2.5 predictions endpoint](https://api.replicate.com/v1/models/stability-ai/stable-audio-2.5/predictions) <br>\n\n\n## Skill Output: <br>\n**Output Type(s):** [guidance, shell commands, configuration, API calls] <br>\n**Output Format:** [Markdown with curl examples and environment-variable setup guidance] <br>\n**Output Parameters:** [1D] <br>\n**Other Properties Related to Output:** [Requires REPLICATE_API_TOKEN for Replicate requests and ffmpeg/ffprobe for mix steps; generated audio is downloaded from the Replicate prediction output.] <br>\n\n## Skill Version(s): <br>\n1.0.7 (source: server release evidence and skill metadata) <br>\n\n## Ethical Considerations: <br>\nUsers should evaluate whether this skill is appropriate for their environment, review any generated or modified files before relying on them, and apply their organization's safety, security, and compliance requirements before deployment. <br>\n\nFile v1.0.7:skill.manifest.json\n\n{\n  \"references\": []\n}\n\nArchive v1.0.6: 11 files, 26079 bytes\n\nFiles: README-INSTALL.md (950b), references/api-credentials.md (3128b), references/audio-post-production.md (9240b), references/generation-diversity.md (25978b), references/generation-quality-checklists.md (10290b), references/random-seed-ritual.md (4011b), references/replicate-api.md (2173b), skill-card.md (2722b), skill.manifest.json (160b), SKILL.md (4556b), _meta.json (135b)\n\nFile v1.0.6:SKILL.md\n\n---\nname: stable-audio-2.5\ndescription: Use when someone wants light instrumental background music — an ambient bed under dialogue or underscore for reels and explainers.\nlicense: MIT\nmetadata:\n  version: \"1.0.6\"\n  provider: replicate\n  replicate_model: stability-ai/stable-audio-2.5\n---\n\n## Shared generation policy\n\n<!-- shared-generation-policy -->\n\nBefore any paid `POST /v1/predictions`:\n\n1. **[Random seed ritual](./references/random-seed-ritual.md)** — always first; derive axes via sum-mod.\n2. **[Generation diversity](./references/generation-diversity.md)** — explicit prompts; rotate ≥2 scenario axes per session.\n3. **[Quality checklists](./references/generation-quality-checklists.md)** — open output files and judge pass/fail before advancing.\n\n# Stable Audio 2.5 (Replicate)\n\nText-to-music for **instrumental background beds** on launch reels. Not a Pruna P-model — runs on [Replicate](https://replicate.com/stability-ai/stable-audio-2.5).\n\n**Mix helper (repo):** [`launch_background_music.py`](https://github.com/PrunaAI/pruna-skills/tree/main/workflows/_shared/scripts/launch_background_music.py) — probes video length, generates bed, mixes under VO with ffmpeg.\n\n## When to use\n\n| Goal | Use this |\n|------|----------|\n| Light instrumental under a concat launch reel | Yes — after final assembly |\n| Under **embedded narration** from scene anchor triple | Yes — mix quiet bed after concat — [audio-post-production.md](./references/audio-post-production.md) |\n| Replace avatar VO | No — bed mixes **under** existing dialogue |\n| Pruna-native audio | No — use [`p-video`](https://github.com/PrunaAI/pruna-skills/tree/main/plugins/p-video/skills/p-video/SKILL.md) audio input instead |\n\n## Environment\n\n```bash\nexport REPLICATE_API_TOKEN=r8_...\n```\n\nRequires **`ffmpeg`** and **`ffprobe`** on PATH for mix step.\n\n## Model input (Replicate)\n\n| Field | Notes |\n|-------|-------|\n| `prompt` | **Required.** Style tags work well — e.g. *Instrumental light electronic pop, soft groove, mellow synth pads, no vocals, 94 BPM* |\n| `duration` | Seconds, 1–190 (match or slightly exceed reel length) |\n| `steps` | 4–8 (default 8) |\n| `cfg_scale` | 1–25 (default 1) |\n| `seed` | Optional integer for reproducibility |\n\nOutput: single **MP3** URL.\n\n## HTTP (curl)\n\n### Create prediction\n\n```bash\ncurl -s -X POST \\\n  -H \"Authorization: Bearer ${REPLICATE_API_TOKEN}\" \\\n  -H \"Content-Type: application/json\" \\\n  -d '{\n    \"input\": {\n      \"prompt\": \"Instrumental light electronic pop bed, soft groove and mellow synth pads, calm positive tech atmosphere, understated background music, no vocals, 94 BPM\",\n      \"duration\": 90,\n      \"steps\": 8,\n      \"cfg_scale\": 1\n    }\n  }' \\\n  \"https://api.replicate.com/v1/models/stability-ai/stable-audio-2.5/predictions\"\n```\n\nPoll `urls.get` until `status` is `succeeded`; download `output` MP3.\n\n## Launch reel integration\n\n### Plan JSON (`background_music`)\n\n```json\n\"background_music\": {\n  \"enabled\": true,\n  \"prompt\": \"Instrumental light electronic pop bed, soft groove and mellow synth pads, calm positive tech atmosphere, understated background music, no vocals, 94 BPM\",\n  \"volume\": 0.12,\n  \"output_name\": \"skills_library_announcement_with_music.mp4\"\n}\n```\n\n### Post-concat bed on any MP4\n\n```bash\npython3 workflows/_shared/scripts/launch_background_music.py \\\n  --video output/my_reel/final.mp4 \\\n  --volume 0.12\n```\n\n## Prompt tips (launch beds)\n\n- Lead with **Instrumental** and **no vocals**\n- Name mood: light, calm, positive, soft groove, mellow — keep *understated* and *background music* so beds sit under VO\n- Optional BPM (**88–98** for tech launch reels; default script uses **94**)\n- Avoid *energetic*, *driving*, or very high BPM — beds should support dialogue, not compete with it\n- Avoid lyrics, song title, or artist name triggers\n- Scene-specific beds (restaurant, jungle, boutique): name the setting but keep groove soft and BPM in the high 80s–mid 90s\n\n## Related\n\n- [audio-post-production.md](./references/audio-post-production.md) — narration + bed layering\n- [gemini-3.1-flash-tts](https://github.com/PrunaAI/pruna-skills/tree/main/plugins/gemini-3.1-flash-tts/skills/gemini-3.1-flash-tts/SKILL.md) — narration voiceover\n- [visual-transition-reel](https://github.com/PrunaAI/pruna-skills/tree/main/plugins/visual-transition-reel/skills/visual-transition-reel/SKILL.md) — concat + optional bed assembly\n- [replicate-api.md](https://github.com/PrunaAI/pruna-skills/tree/main/references/policies/replicate-api.md) — shared Replicate patterns\n\nFile v1.0.6:_meta.json\n\n{\n  \"ownerId\": \"kn7cagwf7q3t0cxrgteb7xk0bh81j0eb\",\n  \"slug\": \"stable-audio-2-5\",\n  \"version\": \"1.0.6\",\n  \"publishedAt\": 1784235618681\n}\n\nFile v1.0.6:references/api-credentials.md\n\n# API credentials (Pruna + Replicate)\n\n**Agent rule:** Before any `POST /v1/predictions`, Replicate prediction, or paid runner — check env vars. If a required key is **missing or empty**, **stop** and tell the user how to sign up. Do not guess, mock, or skip with placeholder keys.\n\n## Pruna P-API\n\n| | |\n|--|--|\n| **Env var** | `PRUNA_API_KEY` |\n| **Header** | `apikey: ${PRUNA_API_KEY}` (not `Authorization: Bearer`) |\n| **Sign up / get key** | [Pruna dashboard](https://dashboard.pruna.ai/) |\n| **Docs** | [Quickstart](https://docs.api.pruna.ai/guides/quickstart) · [pruna-api.md](https://github.com/PrunaAI/pruna-skills/tree/main/workflows/pruna-api/SKILL.md) |\n\n**Used by:** all `p-image*`, `p-video*` tool skills and Pruna workflow runners.\n\n### If `PRUNA_API_KEY` is missing — agent message template\n\n> Pruna generation needs an API key. Sign up or sign in at **[dashboard.pruna.ai](https://dashboard.pruna.ai/)**, create an API key, then set:\n>\n> ```bash\n> export PRUNA_API_KEY=\"your_key_here\"\n> ```\n>\n> Add that to your shell profile or project `.env` (never commit the key). Reply when it’s set and we can continue.\n\n## Replicate\n\n| | |\n|--|--|\n| **Env var** | `REPLICATE_API_TOKEN` |\n| **Header** | `Authorization: Bearer ${REPLICATE_API_TOKEN}` |\n| **Sign up / get token** | [Replicate API tokens](https://replicate.com/account/api-tokens) ([sign in](https://replicate.com/signin) first if needed) |\n| **Docs** | [replicate-api.md](./replicate-api.md) |\n\n**Used by:** `music-2.5`, `gemini-3.1-flash-tts`, `stable-audio-2.5`, `whisperx`, and workflow beds/TTS/song phases.\n\n### If `REPLICATE_API_TOKEN` is missing — agent message template\n\n> This step uses Replicate (song, TTS, transcription, or background bed). Create a token at **[replicate.com/account/api-tokens](https://replicate.com/account/api-tokens)**, then set:\n>\n> ```bash\n> export REPLICATE_API_TOKEN=\"r8_...\"\n> ```\n>\n> Reply when it’s set and we can continue.\n\n## Which key does this job need?\n\n| Task | Keys required |\n|------|----------------|\n| `p-image`, `p-image-edit`, `p-image-upscale`, `p-image-try-on` | `PRUNA_API_KEY` |\n| `p-video`, `p-video-avatar`, `p-video-animate`, `p-video-replace` | `PRUNA_API_KEY` |\n| Music 2.5 song generation | `REPLICATE_API_TOKEN` |\n| Gemini TTS narration | `REPLICATE_API_TOKEN` |\n| Stable Audio background bed | `REPLICATE_API_TOKEN` |\n| WhisperX transcription | `REPLICATE_API_TOKEN` |\n| Music video / explainer (full pipeline) | **Both** — Pruna for stills/video; Replicate for song/TTS/bed as needed |\n\nWhen only one key is missing, suggest **only** that provider’s signup link — not both.\n\n## Security\n\n- Never print full keys in chat or commit them to git.\n- `.env` is gitignored; prefer env vars over hardcoding in plans or manifests.\n- Never embed keys in prompts, manifests, plan JSON, logs, or **subagent task text**.\n- Prefer the **parent agent** to own API calls; do not fan credentials across parallel subagents unless the host documents isolated secret injection.\n- Full rules: [agent-safety.md](https://github.com/PrunaAI/pruna-skills/tree/main/workflows/agent-safety/SKILL.md).\n\nFile v1.0.6:references/audio-post-production.md\n\n# Audio post-production (Pruna + Replicate)\n\nHow to choose and **layer** audio when building reels, multi-scene films, and launch videos.\n\n**Multi-scene narrated films:** use the [scene anchor triple](https://github.com/PrunaAI/pruna-skills/tree/main/workflows/scene-anchor-triple/SKILL.md) — pass TTS to **`p-video`** as `input.audio` with `image` + `last_frame_image`; do not post-mux unless re-render is impossible.\n\n**Visual-only transitions (no VO):** use the [scene anchor pair](https://github.com/PrunaAI/pruna-skills/tree/main/workflows/scene-anchor-pair/SKILL.md) — `duration` instead of `audio`; see [visual-transition-reel](https://github.com/PrunaAI/pruna-skills/tree/main/plugins/visual-transition-reel/skills/visual-transition-reel/SKILL.md).\n\n## Audio-led `p-video` (required when VO/narration exists)\n\nWhen narration, TTS, or a timed audio slice is available **before** video render:\n\n1. Upload the audio file to Pruna (`POST /v1/files`).\n2. Pass `urls.get` as **`input.audio`** on **`p-video`** (or **`p-video-avatar`** for human lip-sync).\n3. **Omit `duration`** — clip length follows the audio (capped at **20s** on P-API); the model syncs motion to speech.\n4. Set **`save_audio`: true** so the full line is embedded in the output clip.\n5. **Probe TTS length** before render — per-scene lines should be **≤ ~19s** or the API truncates the tail even when `audio` is set.\n5. **Concat** clips in order (narration already on each clip). Optional bed mixed **under** VO in post.\n\n**Never** generate silent `p-video` and ffmpeg-mux narration afterward unless re-render is impossible — post-mux **truncates** lines longer than the video slot (common with Gemini TTS).\n\n**Over 20s?** Shorten scene copy → tighten TTS pace in `style_prompt` → split into two scene rows (each with its own triple). See [narrated-multi-scene](https://github.com/PrunaAI/pruna-skills/tree/main/plugins/narrated-multi-scene/skills/narrated-multi-scene/SKILL.md) duration gate.\n\nHelper: [`p_video_payload.py`](https://github.com/PrunaAI/pruna-skills/tree/main/workflows/_shared/scripts/p_video_payload.py) — `build_p_video_payload(...)` enforces omitting `duration` when `audio_url` is set.\n\n| Workflow | Audio source | `p-video` fields |\n|----------|--------------|------------------|\n| Dog plush / story film | Gemini TTS per scene | `image` + `last_frame_image` + `audio` |\n| Music video performance | Song slice per cut | `image` + `audio` |\n| Music video B-roll | Song slice (optional) | `image` + `audio` or `duration` only |\n| Viking narrator beats | Gemini TTS | `image` + `last_frame_image` + `audio` |\n\n## Tool picker\n\n| Need | Tool | Skill |\n|------|------|-------|\n| Cinematic clip with model-generated sound | `p-video` (`save_audio`, optional uploaded `audio`) | [p-video](https://github.com/PrunaAI/pruna-skills/tree/main/plugins/p-video/skills/p-video/SKILL.md) |\n| Lip-sync / duration locked to VO | Upload audio → `p-video` with `audio` | [p-video](https://github.com/PrunaAI/pruna-skills/tree/main/plugins/p-video/skills/p-video/SKILL.md) |\n| Documentary / story narrator | [Gemini 3.1 Flash TTS](https://replicate.com/google/gemini-3.1-flash-tts) | [gemini-3.1-flash-tts](https://github.com/PrunaAI/pruna-skills/tree/main/plugins/gemini-3.1-flash-tts/skills/gemini-3.1-flash-tts/SKILL.md) |\n| Light instrumental under dialogue | [Stable Audio 2.5](https://replicate.com/stability-ai/stable-audio-2.5) | [stable-audio-2.5](../SKILL.md) |\n| Full song with sung vocals | [Music 2.5](https://replicate.com/minimax/music-2.5) | [music-2.5](https://github.com/PrunaAI/pruna-skills/tree/main/plugins/music-2.5/skills/music-2.5/SKILL.md) |\n| Speaking on-camera character | `p-video-avatar` | [p-video-avatar](https://github.com/PrunaAI/pruna-skills/tree/main/plugins/p-video-avatar/skills/p-video-avatar/SKILL.md) |\n\n**Env:** Pruna calls need `PRUNA_API_KEY`; Replicate audio tools need `REPLICATE_API_TOKEN`. Assembly steps need **`ffmpeg`** / **`ffprobe`**.\n\n## Layering matrix\n\n| Stack | Primary audio | Secondary | Mix notes |\n|-------|---------------|-----------|-----------|\n| **Silent B-roll** | — | — | Concat video only |\n| **Native `p-video` sound** | Model output | — | Keep `save_audio` default; normalize in assembly if scenes differ |\n| **Narration only (fallback)** | Gemini TTS | — | Post-mux only when audio-led `p-video` is not suitable — prefer **Pipeline B** below |\n| **Bed only** | Stable Audio bed | — | [`launch_background_music.py`](https://github.com/PrunaAI/pruna-skills/tree/main/workflows/_shared/scripts/launch_background_music.py) |\n| **Narration + bed (preferred)** | Gemini TTS → **`p-video` `audio`** | Stable Audio (quiet) | TTS uploaded to Pruna drives clip length + sync; bed mixed in post under narration (~0.08–0.15) |\n| **Avatar VO + bed** | `p-video-avatar` dialogue | Stable Audio bed | Same bed pattern as replace/launch reels — bed **under** existing speech |\n| **Music video** | Music 2.5 full song | — | [music-video](https://github.com/PrunaAI/pruna-skills/tree/main/plugins/music-video/skills/music-video/SKILL.md) |\n\n## Recommended pipelines\n\n### A — Narrated multi-scene B-roll (**preferred — scene anchor triple**)\n\n```text\nPhase 0 — intake: scene table with start/end still prompts + narration lines\nPhase 1 — hero + p-image-edit start stills + end stills (parallel)\nPhase 2 — Gemini TTS per scene (parallel) → upload each to /v1/files\nPhase 3 — p-video per scene: input.image + input.last_frame_image + input.audio (parallel; omit duration)\nPhase 4 — ffmpeg concat (VO embedded; frame chain via shared end/start URLs)\nPhase 5 — optional Stable Audio bed under narration\n```\n\n**Scene anchor triple:** same pattern as first/last frame pairing — `audio` is the third required upload per scene row. **`p-video-avatar`:** portrait + optional `last_frame_image` + uploaded `audio`.\n\n### A′ — Post-mux narration (fallback only)\n\nUse only when you already have silent clips and cannot re-render. Risk: TTS longer than clip slots → cut-off VO.\n\n```text\nPhase 3 — p-video I2V without audio → concat → mux TTS in ffmpeg\n```\n\n### C — Launch / product reel (existing pattern)\n\n```text\nPhase 1 — p-video-avatar or replace reel → concat\nPhase 2 — Stable Audio bed via launch_background_music.py (bed under VO, not replacing it)\n```\n\n## ffmpeg mixing (conceptual)\n\n**Narration onto silent concat** (single VO file):\n\n```bash\nffmpeg -y -i concat_video.mp4 -i narration.mp3 \\\n  -map 0:v -map 1:a -c:v copy -c:a aac -b:a 192k -shortest output_with_vo.mp4\n```\n\n**Bed under existing narration + video** (same pattern as `launch_background_music.py`):\n\n```text\n[1:a]volume=0.12,aloop=...[bed];\n[0:a][bed]amix=inputs=2:duration=first[aout]\n```\n\nNarration / avatar dialogue stays on stream `0:a`; bed is stream `1:a` at low volume.\n\n**Bed on silent concat** — loop a short generated clip to full video length (no per-assemble Stable Audio call):\n\n```text\n[1:a]volume=0.12,aloop=loop=-1:size=2e+09[bed]  →  map video + [bed], -shortest\n```\n\nPlan field `\"reuse_bed\": true` skips regeneration when `audio/launch_bed.mp3` exists. Delete that file (or set `reuse_bed: false`) only when you want a new prompt or seed.\n\n## Intake questions (audio)\n\nAsk before generating paid audio or video:\n\n| Topic | Questions |\n|-------|-----------|\n| **Primary voice** | Narrator (Gemini TTS), on-screen avatar (`p-video-avatar`), or native `p-video` sound only? |\n| **Narration scope** | Per-scene lines vs one continuous VO track? |\n| **Music / bed** | None, instrumental bed only, or full song ([music-2.5](https://github.com/PrunaAI/pruna-skills/tree/main/plugins/music-2.5/skills/music-2.5/SKILL.md))? |\n| **Sync strategy** | **Preferred:** TTS → Pruna upload → **`p-video` / `p-video-avatar` with `audio`** (clip length = audio). Post-mux only as fallback. |\n| **Levels** | Bed volume target (default ~0.12 under avatar VO; ~0.08–0.12 under Gemini narration)? |\n\n## Manifest fields\n\n```json\n{\n  \"narration\": { \"enabled\": true, \"voice\": \"Sulafat\", \"mode\": \"per_scene\" },\n  \"background_music\": { \"enabled\": true, \"reuse_bed\": true, \"volume\": 0.10, \"prompt\": \"Instrumental ... no vocals\" },\n  \"p_video_audio\": { \"save_audio\": true }\n}\n```\n\n## Limitations (from [P-Video on Replicate](https://replicate.com/prunaai/p-video))\n\n- Native SFX/dialogue quality varies — for premium voice realism, prefer **Gemini TTS** or **`p-video-avatar`**, then optionally mix a bed.\n- Multi-speaker native audio can drift; dedicated TTS per role is safer for narration-heavy cuts.\n- Extreme camera motion and complex multi-scene stories are weaker than **frame-anchored chaining** + per-scene prompts — see [p-video](https://github.com/PrunaAI/pruna-skills/tree/main/plugins/p-video/skills/p-video/SKILL.md) **First / last frame chaining**.\n\n## Related\n\n- [parallel-execution.md](https://github.com/PrunaAI/pruna-skills/tree/main/policies/parallel-execution.md) — phased vs parallel when frames chain\n- [narrated-multi-scene](https://github.com/PrunaAI/pruna-skills/tree/main/plugins/narrated-multi-scene/skills/narrated-multi-scene/SKILL.md)\n- [pruna-generative-pipeline](https://github.com/PrunaAI/pruna-skills/tree/main/docs/WORKFLOW-RECIPES.md)\n\nFile v1.0.6:references/generation-diversity.md\n\n# Generation diversity (all models)\n\nOne checklist so **every** Pruna output — **`p-image`**, **`p-video`**, try-on, avatar, replace, animate — is as **diverse** as the brief allows. Details live in linked docs; this page is the agent shortcut.\n\nUse the **full** checklist here for every generation.\n\n## Contents\n\n- [Three steps (every job)](#three-steps-every-job)\n- [Explicit prompt structure](#explicit-prompt-structure-required)\n- [Text & typography by model](#text--typography-by-model)\n- [SSoT axis derivation](#ssot-axis-derivation-sum-mod)\n- [Scenario axes](#scenario-axes-rotate-across-outputs)\n- [Render categories](#render-categories)\n- [Crowded scenes](#crowded-scenes-p-image)\n- [Body type spread](#body-type-spread)\n- [Location-matched crowds](#location-matched-crowds)\n- [Group classes](#group-classes--courses)\n- [Framing & camera](#framing--camera)\n- [Scene spice](#scene-spice-when-it-fits)\n- [Photoreal anti-slop](#photoreal-anti-slop-neon--stylized-briefs)\n- [Aspect ratio](#aspect-ratio-multi-example-sets)\n- [By model](#by-model-minimum-diversity)\n- [When not to maximize diversity](#when-not-to-maximize-diversity)\n- [Anti-patterns](#anti-patterns)\n\n## Three steps (every job)\n\n1. **[Random seed ritual](./random-seed-ritual.md) (SSoT)** — **always first**, before the prompt. Generate a fresh random string, **state it in the turn**, derive axes via [sum-mod](#ssot-axis-derivation-sum-mod). **Do not** pass the ritual string to API `seed`. **One new ritual string per independent generation**; reuse only on same-brief slop retry.\n2. **Write an [explicit prompt](#explicit-prompt-structure-required)** — name specific people, animals, objects, actions, setting, and camera/light. Add text/typography only when the brief needs it — see [text rules by model](#text--typography-by-model).\n3. **Diversify the scenario row** — change at least **two axes** from the previous output in the same session (cast, setting, camera, **`render_category_tag`**, **aspect_ratio**, creatures, props, … — unless user asked for continuity).\n4. **Log** — `ritual_seed`, axes chosen, prediction id (manifest or turn text).\n\n## Explicit prompt structure (required)\n\n**Vague prompts produce generic AI slop.** After the ritual and axis picks, every still prompt must be **specific and dynamic** — concrete nouns, frozen actions, named places. Prefer playground/creative briefs over marketing abstractions.\n\n**Name at least four of these per prompt (log tags in manifest):**\n\n| Clause | Log as | Agent must specify |\n|--------|--------|-------------------|\n| **People** | `cast_descriptor` | Named role + age band + expression (`fearless grandmother in floral apron`, not `woman`) |\n| **Animals / creatures** | `creature_tag` | Species + attitude (`otter DJ`, `luna moth knight`, `VIP anglerfish`) |\n| **Objects** | `prop_tag` | Concrete props (`vinyl record`, `chrome rocket sled`, `velvet rope`, `tiny boombox`) |\n| **Action** | `action_tag` | Frozen mid-motion verb (`scratching vinyl`, `lassoing runaway taco truck`, `cape mid-swing`) |\n| **Duration** | `duration_tag` | When timing matters (`1970s`, `8PM`, `45-minute spin class`, `Saturday-morning cartoon`) |\n| **Setting** | `setting_tag` | Named place + era + materials (`packed 1970s roller rink`, `abyss-depth jellyfish nightclub`, `Monument Valley dust storm`) |\n| **Text / typography** | `text_spec` | Only when brief needs readable type — exact strings + surface (see [by model](#text--typography-by-model)) |\n| **Camera + light** | `camera_tag`, `lighting_tag` | `fish-eye lens`, `tilt-shift macro`, `teal-magenta cinematic`, `golden hour sparkle` |\n| **Style** | `render_category_tag` | Medium (`cel-shaded anime`, `baroque oil painting`, `ink-wash storybook`, `photoreal documentary`) |\n\n**Template:**\n\n```text\n{people and/or creatures} {action} with/at {specific objects} in {named setting},\n{style or era cues}, {camera_tag}, {lighting_tag}\n```\n\n**Good examples (dynamic / specific):**\n\n```text\nDisco ball reflections on an otter DJ scratching vinyl at a packed 1970s roller rink,\nfish-eye lens, glitter confetti mid-air, funky energy\n```\n\n```text\nBioluminescent jellyfish nightclub at abyss depth, VIP anglerfish in sunglasses at velvet rope,\nteal-magenta cinematic lighting\n```\n\n```text\nCorgi cowboy lassoing a runaway taco truck through Monument Valley dust storm,\npulp western poster energy, dynamic diagonal composition\n```\n\n**Anti-pattern:** `cool cyberpunk portrait, neon vibes` — no subject, no action, no place. **Right:** name who, what they're doing, where, with which props.\n\n## Text & typography by model\n\n**Never use negation to suppress text** — `no text`, `without signs`, `no typography` often **invoke** the thing you are trying to avoid. Describe surfaces positively when you want blank walls (`plain unmarked walls`, `matte unprinted props`).\n\n| Model | Prompt upsampling | Typography in prompt |\n|-------|-------------------|----------------------|\n| **`p-image`** | **No** effective prompt upsampling | **Avoid** dense readable-type requests unless user explicitly wants `text_rendering`. Short prompts; skip `readable`, `legible`, `headline`, multi-sign lists — they drift to gibberish. Collage triggers still apply: [interactive-explainer-prompts.md](https://github.com/PrunaAI/pruna-skills/tree/main/workflows/interactive-explainer-prompts.md) (`flat lay`, `grid`, `collage`, …). |\n\n**`p-image` text hygiene:** prefer scenes without copy. If a screen appears: `monitor soft colorful blur glow only` — not legible UI unless the user explicitly asked for readable text (then simplify the brief or drop copy).\n\n**Collage triggers (all T2I models):** still avoid `flat lay`, `packshot`, `grid`, `collage`, `montage`, `contact sheet`, `split`, `before and after` — use `single frame`, `one camera angle` instead. Full table: [interactive-explainer-prompts.md](https://github.com/PrunaAI/pruna-skills/tree/main/workflows/interactive-explainer-prompts.md).\n\n## SSoT axis derivation (sum-mod)\n\nAfter stating `ritual_seed` (random string), derive prompt choices — sum Unicode/ASCII char codes, mod list length:\n\n```text\nRATIOS = [\"1:1\", \"16:9\", \"9:16\", \"4:3\", \"3:4\", \"3:2\", \"2:3\"]\naspect_ratio  ← RATIOS[ sum(codes(ritual_seed)) % 7 ]\ncamera_tag    ← camera_tags[ sum(codes(ritual_seed[0:4])) % len(camera_tags) ]\nrender_tag    ← render_tags[ sum(codes(ritual_seed[4:8])) % len(render_tags) ]\n```\n\n`camera_tags` and `render_tags` — see [framing & camera](#framing--camera) and [render categories](#render-categories). State derived picks in the turn (*\"Aspect ratio: 16:9, camera: over-shoulder\"*).\n\n**User `api_seed`:** when the user supplies an integer for reproducibility, pass it as `input.seed` — separate from the ritual string.\n\n## Scenario axes (rotate across outputs)\n\n| Axis | Vary with | Applies to |\n|------|-----------|------------|\n| **Cast** | age, ethnicity, gender, archetype, **hairstyle**, **body type** (rotate — see [below](#body-type-spread)), disability aids (wheelchair, cane), visible age band twice in prompt | all person/content gens |\n| **Medium** | `render_category_tag` — rotate across [render categories](#render-categories) | `p-image`, avatar stills |\n| **Setting** | unique `setting_tag` — specific room/street/venue/era, not repeat adjacent rows | stills + video plates |\n| **Camera** | `camera_tag` — rotate across [framing ladder](#framing--camera); never default MC facing lens | stills, `video_prompt` |\n| **Lighting** | `lighting_tag` — golden hour · neon · overcast · practical | stills, video mood |\n| **Motion** | unique `video_prompt` per clip | `p-video`, `p-video-avatar`, animate |\n| **Voice** | natural `voice_script`; one `voice` preset per character | avatar, TTS-led video |\n| **Seed** | new ritual string per **independent** job; reuse only on same-brief slop retry | all generation skills |\n| **Aspect ratio** | different `aspect_ratio` per independent still in a batch — see [below](#aspect-ratio-multi-example-sets) | `p-image`, `p-image-edit` |\n| **Crowd density** | layered background population + activity cues — see [below](#crowded-scenes-p-image) | `p-image` plates with busy worlds |\n\nFull style/camera/lighting ladders: [visual-variety-bible.md](https://github.com/PrunaAI/pruna-skills/tree/main/shared/visual-variety-bible.md). Persona + try-on bar: [realistic-persona-showcase.md](https://github.com/PrunaAI/pruna-skills/tree/main/shared/realistic-persona-showcase.md).\n\n## Render categories\n\nRotate **`render_category_tag`** (and log it) so diversity batches cover more than photoreal portraits or anime. Category families below mirror arena leaderboards — pick a **different tag per independent output**.\n\n**Random seed ritual still applies** to every generation in [step 1](#three-steps-every-job); categories describe *what* to vary, not *when* to pick `seed`.\n\n### Text-to-image — `p-image`\n\nSources: [Arena text-to-image](https://arena.ai/leaderboard/text-to-image) · [AA text-to-image](https://artificialanalysis.ai/image/leaderboard/text-to-image)\n\n**Unified `render_category_tag`** (Arena bucket = tag — pick one per still):\n\n`product_branding_commercial` · `3d_imaging_modeling` · `cartoon_anime_fantasy` · `photoreal_cinematic` · `art` · `portraits` · `nature_environment` · `animals_creature` · `text_rendering`\n\n| Tag | Typical prompt lane |\n|-----|---------------------|\n| `product_branding_commercial` | single product on seamless studio, person + product in named setting, showroom (not `flat lay` / `packshot` words) |\n| `3d_imaging_modeling` | CG film still, clay/stop-motion, rounded 3D forms |\n| `cartoon_anime_fantasy` | cel anime, fantasy character, crowded stylized world |\n| `photoreal_cinematic` | documentary crowd scenes, film-scale wide, urban march |\n| `art` | oil, watercolor, gouache, charcoal, flat vector |\n| `portraits` | single-subject editorial or documentary portrait (crowd optional behind) |\n| `nature_environment` | landscape-wide; subject small in frame |\n| `animals_creature` | named species + handler; crowded market/park when it fits |\n| `text_rendering` | **user-requested only** — otherwise no readable text |\n\nLog `render_category_tag` in manifest. Combine with [crowded scenes](#crowded-scenes-p-image), [body type](#body-type-spread), and [scene spice](#scene-spice-when-it-fits) when the brief allows.\n\n### Image edit — `p-image-edit`\n\nSources: [Arena image edit](https://arena.ai/leaderboard/image-edit) · [AA image editing](https://artificialanalysis.ai/image/leaderboard/editing)\n\nArena modalities: `single_image_edit` · `multi_image_edit`\n\nEdit diversity tags: `background_swap` · `relight` · `wardrobe_on_plate` · `pose_or_angle_delta` · `multi_ref_composite` · `region_inpaint`\n\nVary **instruction** and **what changes** while identity URL stays fixed on character arcs.\n\n### Text-to-video — `p-video`\n\nSources: [Arena text-to-video](https://arena.ai/leaderboard/text-to-video) · [AA text-to-video](https://artificialanalysis.ai/video/leaderboard/text-to-video)\n\nMotion/scene tags: `character_performance` · `landscape_broll` · `urban_street` · `product_demo` · `abstract_mood` · `crowd_scene` · `dialogue_beat`\n\nRotate `video_prompt` grammar, start plate world, and `camera_tag` per clip.\n\n### Image-to-video — `p-video` (+ plate upload)\n\nSources: [Arena image-to-video](https://arena.ai/leaderboard/image-to-video) · [AA image-to-video](https://artificialanalysis.ai/video/leaderboard/image-to-video)\n\nPlate-driven tags: `animate_hero_still` · `camera_move_on_plate` · `environmental_parallax` · `avatar_lip_sync` · `hands_or_prop_motion`\n\nMatch motion to what the **still** already shows — do not contradict the plate.\n\n### Video edit — `p-video-replace` (and edit-style video)\n\nSource: [Arena video edit](https://arena.ai/leaderboard/video-edit)\n\nEdit tags: `face_recast` · `wardrobe_swap` · `accessory_swap` · `background_replace` · `object_in_hand_swap` · `style_transfer_on_subject`\n\nSame-gender / identity rules for talking-head beats still apply — see [visual-variety-bible.md](https://github.com/PrunaAI/pruna-skills/tree/main/shared/visual-variety-bible.md).\n\n## Crowded scenes (`p-image`)\n\nWhen the brief asks for **busy**, **crowded**, or **lively** worlds — not a lone subject on a blank wall — stack density in the prompt:\n\n1. **Three depth layers** — sharp foreground subject · readable midground faces/hands/props · landmark bokeh (stage, temple, billboards, ferris wheel).\n2. **Named population count** — `hundreds of pedestrians`, `dozens of faces in midground`, `20+ tiny clay figures` (stylized sets need explicit counts; models under-deliver on vague \"busy\").\n3. **Activity verbs** — raised hands, umbrellas open, food steam, confetti, market haggling, commuters pressed shoulder-to-shoulder.\n4. **Shallow DOF + single subject** — `single subject one frame` keeps one identity readable while the crowd stays behind them.\n5. **Age & angle lock** — repeat age band twice (`woman in her late 50s, visibly fifty`) and use [framing & camera](#framing--camera) — models drift younger, center-frame, and front-facing without it.\n\n| Crowd family | Density cues |\n|--------------|--------------|\n| **Urban rush** | crosswalk stripes, wet reflections, umbrellas, billboard bokeh |\n| **Festival / parade** | confetti, raised hands, costume layers, smoke haze |\n| **Market / bazaar** | overflowing stalls, hanging goods, steam, price tags as color blobs |\n| **Transit crush** | strap hangers, door windows, blurred faces pressed together |\n| **Stylized miniature** | counted clay/figurine shoppers (`20+`), cramped aisle, stacked crates |\n| **Institutional / ER** | framed oil portraits on beige walls, triage number board, wall sanitizer, vending machine, scuffed linoleum, TV blur, mixed-age seated patients |\n| **Urban march / protest** | named city, local landmarks, multiracial crowd cues separate from hero — see [location-matched crowds](#location-matched-crowds) |\n| **Group fitness class** | class name + duration, mixed-gender riders, realistic warm studio light — see [group classes](#group-classes--courses) |\n\n**Anti-pattern:** one blurred smear behind a portrait — name **what** the crowd is doing and **where** layers sit. **Institutional** scenes (ER, airport, classroom) need `benches full`, `standing room only`, or `shoulder-to-shoulder` — otherwise models default to a quiet hallway. Name **set dressing** too: framed portraits on walls, triage number board, vending machine glow, scuffed linoleum — generic mint corridors read AI-empty.\n\n## Body type spread\n\nModels default to one “average fitness” body. In diversity batches, **name build on the hero and vary background bodies**:\n\n| Build tag | Prompt cue |\n|-----------|------------|\n| **Plus-size / curvy** | `plus-size`, `curvy build`, `full-figured` |\n| **Athletic / muscular** | `broad shoulders`, `muscular arms`, `athletic build` |\n| **Petite / slim** | `petite frame`, `slim build`, `narrow shoulders` |\n| **Tall / lanky** | `tall and lanky`, `6-foot frame`, `long limbs` |\n| **Stocky / heavyset** | `stocky build`, `heavyset`, `barrel chest` |\n| **Lean wiry** | `lean wiry frame`, `weathered thin face` |\n\n**Rule:** rotate build across independent panels in a session — not every hero “athletic build”. Background crowd should mix ages **and** silhouettes (`elderly thin woman`, `heavyset man`, `pregnant woman seated`, `toddler on lap`).\n\n## Location-matched crowds\n\nWhen the prompt names a **real city or country**, background faces must match that place’s **demographic mix** — not clone the hero’s ethnicity.\n\n| Wrong | Right |\n|-------|--------|\n| South Asian hero + only South Asian protesters in “New York” | Hero is one identity; crowd explicitly `multiracial NYC march — Black, Latino, white, East Asian protesters` |\n| “Dense city march” with no geography | Name city + 3–4 crowd ethnicity cues + local landmarks (yellow cabs, art deco towers, steam vent) |\n| Festival in Lagos with only Nordic faces | Match crowd to `setting_tag` region |\n\n**Prompt pattern:** lock hero cast in sentence 1; sentence 2 lists **four+ distinct background silhouettes** unrelated to hero ethnicity; sentence 3 names **local landmarks** so the plate cannot read as generic stock.\n\n**Applies to:** protests, airports, transit, street markets, sports crowds — any scene where “crowded” implies a real place.\n\n## Group classes & courses\n\nWhen the scene is a **class, workshop, or team activity**, name the **course type** and **who else is in the room** — models default to monochrome crowds (all men, all one age).\n\n| Specify | Example cues |\n|---------|----------------|\n| **Class type** | `45-minute evening spin class`, `beginner yoga flow`, `HIIT bootcamp circuit` |\n| **Room realism** | warm overhead track lights, mirror wall, rubber floor, water bottles, towels — **not** magenta-cyan neon strips unless brief is explicitly nightclub |\n| **Gender mix** | hero is one person; crowd `mixed-gender class — women with ponytails, men with beards, nonbinary cyclist` |\n| **Body + age mix** | plus-size rider, petite woman, athletic man, woman in her 50s — same as [body type spread](#body-type-spread) |\n\n**Lighting rule for fitness:** real boutique studios are **dim warm overhead** or **single spotlight on instructor** — avoid `split gel`, `neon LED strips`, `magenta-cyan` on photoreal gym plates; those read AI-fake.\n\n**Prompt pattern:** `Documentary fitness portrait` + class name + instructor on bike at front + `20+ mixed-gender cyclists` with 3–4 named background silhouettes + realistic room props.\n\n## Framing & camera\n\nModels default to **centered subject, eyes at camera**. In diversity batches, **rotate `camera_tag` and frame placement** every row — log both in manifest.\n\n**Gaze rule:** `glance off-lens`, `profile`, `back to camera`, `looking down at [prop]`, or `watching the crowd` — **not** `facing camera` or `looking at viewer` unless the user asked for a direct-address avatar plate.\n\n**Placement rule:** name where the subject sits in frame — `left third`, `right third`, `lower right corner`, `edge of frame`, `small in environmental wide` — **not** centered mugshot every time.\n\n| `camera_tag` | Prompt cue |\n|--------------|------------|\n| **Overhead / bird's eye** | `overhead aerial view`, `top-down`, `drone shot looking straight down` |\n| **High corner** | `high angle from corner`, `surveillance-style downward angle` |\n| **Worm's eye** | `ground-level worm's eye`, `camera on pavement` |\n| **Crane-down** | `slight high angle crane-down` |\n| **Over-shoulder** | `over-shoulder from behind`, `seen past someone's shoulder` |\n| **Profile / side** | `profile side angle`, `walking across frame` |\n| **From behind** | `back to camera`, `three-quarter from behind` |\n| **Dutch tilt** | `dutch tilt` — tension scenes only |\n| **Through crowd** | `subject visible through gap in crowd`, `foreground heads out of focus` |\n\n**Batch rule:** no two adjacent stills share the same `camera_tag` **and** placement corner (e.g. don't do `left third` twice in a row).\n\nAvatar / lip-sync exception: face must stay readable and mouth visible — use `slight angle from the side` or `three-quarter`, still **off-center** and **off-lens gaze** when not delivering VO to camera.\n\n## Scene spice (when it fits)\n\nDefault plates are person + crowd + place. Add **one or two specific attributes** when the setting naturally supports them — not random clutter on every row.\n\n| Spice type | When to add | Example |\n|------------|-------------|---------|\n| **Animals** | setting implies them | dog park → `golden retriever on leash`; harbor → `seagulls overhead`; rooftop → `pigeons on water tower`; parade → `police horse midground` |\n| **Held / worn props** | role or weather | `red umbrella tucked under arm`, `wire beekeeper smoker`, `chipped ceramic mug`, `sample strawberry basket` |\n| **Micro-detail** | one thumb-stopping oddity | `muddy paw prints on pavement`, `honey jar on crate`, `green parade beads on fence` |\n\nCamera and placement live in [framing & camera](#framing--camera) — not optional spice.\n\n**Rule:** pick **at most two** spice items per prompt. They must answer “what would a photographer notice here?” — not a checklist dump.\n\n**Skip spice when:** product hero, avatar MC talking head, try-on full-body (garment is the focus), or minimal studio brief.\n\n## Photoreal anti-slop (neon / stylized briefs)\n\nStylized settings still need **documentary skin discipline** or outputs go waxy:\n\n- Lead with `documentary portrait, natural skin pores, not CGI, not illustration` even for neon/cyberpunk worlds.\n- Prefer **worn real materials** — matte leather, faded denim, scratched CRT bezels, sticky carpet — over `holographic puffer`, `chrome armor`, `HUD`.\n- Name **gritty location cues** — basement arcade, wet alley, scuffed linoleum — not abstract `neon corridor`.\n- Background crowd faces need **imperfect texture**; blur is fine, plastic skin in midground is not.\n\n## Aspect ratio (multi-example sets)\n\nWhen generating **two or more** stills in one session (playground grid, demo batch, mood board), give each independent output a **different** `aspect_ratio` unless the user locked a format.\n\n**Allowed `p-image` values:** `1:1` · `16:9` · `9:16` · `4:3` · `3:4` · `3:2` · `2:3`\n\n**How to pick:** after the [random seed ritual](./random-seed-ritual.md), use [sum-mod](#ssot-axis-derivation-sum-mod) on `ritual_seed` — state it in the turn (*\"Aspect ratio: 16:9\"*). Do **not** default every example to `9:16` or `1:1`.\n\n| Ratio | Typical use |\n|-------|-------------|\n| `9:16` | vertical UGC, full-body fashion, avatar talking head |\n| `16:9` | environmental wide, cinematic landscape plate |\n| `3:4` | editorial portrait, try-on full-body |\n| `4:3` | classic portrait, product + person |\n| `1:1` | packshot grid, social tile |\n| `3:2` · `2:3` | magazine / poster crops |\n\nMatch prompt framing to ratio (e.g. `16:9 horizontal wide shot`, `9:16 vertical full body`). **`p-image-try-on`** inherits plate size when `preserve_input_size: true` — diversify person plates first.\n\n**Same character arc:** one ratio for the whole chain unless the user asks for reframes.\n\n## By model (minimum diversity)\n\n| Model | Besides ritual seed, always vary |\n|-------|-----------------------------------|\n| **`p-image`** | cast/creature + objects + action + setting + camera + **`render_category_tag`** + **aspect_ratio**; [explicit structure](#explicit-prompt-structure-required); [text hygiene](#text--typography-by-model) (no upsampling) |\n| **`p-image-edit`** | edit tag + setting/angle delta; same identity URL |\n| **`p-image-try-on`** | person plate world + garment complexity; preserve scene |\n| **`p-image-upscale`** | N/A on prompt — diversify **source** stills |\n| **`p-video`** | motion/scene tag + `video_prompt`; differ start plates per scene |\n| **`p-video-avatar`** | `video_prompt` + still world per scene; lock voice per character |\n| **`p-video-animate`** | persona still style/setting per slider ref |\n| **`p-video-replace`** | video-edit tag + full cast spread on showcase reels |\n\n## When **not** to maximize diversity\n\n- **Same character arc** — lock hero plate URL, one `voice`, cast descriptor; vary only setting/angle/motion per scene.\n- **User asked for continuity** — match their cast and approved plates.\n- **Draft → final** — same prompt; change only `draft: false`. Use `api_seed` only if user locked API reproducibility.\n\n## Anti-patterns\n\n| Wrong | Right |\n|-------|--------|\n| Copy doc example ritual strings | [Random seed ritual](./random-seed-ritual.md) — fresh string each time |\n| Pass ritual string as API `seed` | Ritual is SSoT planning only; `api_seed` when user requests |\n| White wall + MC CU on every demo | Rotate setting + camera + cast |\n| One `video_prompt` for whole reel | Unique motion per scene row |\n| New ritual string mid avatar chain on same brief | Reuse `ritual_seed` until recast or new independent output |\n| Same aspect ratio on every playground example | Rotate `1:1` · `16:9` · `9:16` · `4:3` · `3:4` · `3:2` · `2:3` per [aspect ratio rules](#aspect-ratio-multi-example-sets) |\n| Every hero same athletic body | Rotate [body type spread](#body-type-spread) |\n| Generic hospital hallway | Named ER set dressing + mixed body types in crowd |\n| `holographic` / `chrome` on photoreal cyber scenes | Worn leather, scratched cabinets, documentary skin cues |\n| Monoculture crowd in a named global city | [Location-matched crowds](#location-matched-crowds) — hero ≠ background ethnicity |\n| Magenta-cyan neon on photoreal gym | Warm overhead studio light, mirror wall, real spin bikes |\n| All-male or all-female group class | [Group classes](#group-classes--courses) — mixed-gender background cues |\n| Centered subject every frame | [Framing & camera](#framing--camera) — rotate `camera_tag` + placement |\n| Subject facing camera / at viewer | Off-lens gaze, profile, from behind, or watching crowd |\n| Random animals with no setting reason | Animals only when place implies them |\n| Every stylized panel is anime | Rotate [render categories](#render-categories) — use `cartoon_anime_fantasy` at most once per batch |\n| Vague `cool portrait, neon vibes` | [Explicit structure](#explicit-prompt-structure-required) — named subject, action, objects, setting |\n| `no text` / `without signage` in prompt | Negation invokes text — use [text rules by model](#text--typography-by-model) |\n| Dense typography on **`p-image`** | Drop copy or simplify the brief — `p-image` has no prompt upsampling |\n\n## Related\n\n- [generation-quality-checklists.md](./generation-quality-checklists.md) — core + model checklists\n- [staged-generation-gate.md](https://github.com/PrunaAI/pruna-skills/tree/main/references/policies/staged-generation-gate.md) — approval phases\n\nFile v1.0.6:references/generation-quality-checklists.md\n\n# Generation quality checklist hub\n\nUse this as the shared quality gate across models and workflows.\nRun the **Core checklist** for every generation job, then run the model-specific checklist.\n\n## Who applies these checklists?\n\n**The coding agent** — by **opening the real output files** (images, video, or audio) and reviewing them with vision. These checklists are **not** automated test scripts. There is no separate scoring service: the agent reads each item and judges pass or fail from what it sees and hears.\n\nTypical flow:\n\n1. **Generate or download** the asset to a local path (`stills/`, `clips/`, etc.).\n2. **Inspect the file** — view the image, watch the video clip, or listen to narration when the checklist covers audio.\n3. Run the **Core checklist** (below), then the **model-specific checklist** for that job.\n4. **If something fails** — note which items failed, adjust prompt / settings / seed, and regenerate **only that asset** (do not advance to expensive video steps on a bad still).\n5. **If it passes** — show the user the file paths (and previews when helpful). In workflows, still follow [staged-generation-gate.md](https://github.com/PrunaAI/pruna-skills/tree/main/references/policies/staged-generation-gate.md): agent checklist review happens **before** you ask the user to approve stills or clips.\n\nThe user's **approve plan / approve stills / approve clips** gates are separate. Agent checklists catch obvious problems early so the user is not asked to sign off on broken outputs.\n\nMaintenance rule: keep tool/workflow mapping only in this file to avoid link drift.\n\n## Match map (tool -> checklist -> workflows)\n\n| Tool/model | Checklist | Common workflows |\n|------------|-----------|---------------|\n| `p-image` | [`p-image-quality-checklist.md`](https://github.com/PrunaAI/pruna-skills/tree/main/image/p-image-quality-checklist.md) · persona bar: [`realistic-persona-showcase.md`](https://github.com/PrunaAI/pruna-skills/tree/main/shared/realistic-persona-showcase.md) | [`image-to-video`](https://github.com/PrunaAI/pruna-skills/tree/main/plugins/image-to-video/skills/image-to-video/SKILL.md), [`narrated-multi-scene`](https://github.com/PrunaAI/pruna-skills/tree/main/plugins/narrated-multi-scene/skills/narrated-multi-scene/SKILL.md), [`WORKFLOW-RECIPES`](https://github.com/PrunaAI/pruna-skills/tree/main/docs/WORKFLOW-RECIPES.md) |\n| `p-image-edit` | [`p-image-edit-quality-checklist.md`](https://github.com/PrunaAI/pruna-skills/tree/main/image/p-image-edit-quality-checklist.md) | [`avatar-single-scene`](https://github.com/PrunaAI/pruna-skills/tree/main/plugins/avatar-single-scene/skills/avatar-single-scene/SKILL.md), [`avatar-multi-scene`](https://github.com/PrunaAI/pruna-skills/tree/main/plugins/avatar-multi-scene/skills/avatar-multi-scene/SKILL.md), [`WORKFLOW-RECIPES`](https://github.com/PrunaAI/pruna-skills/tree/main/docs/WORKFLOW-RECIPES.md) |\n| `p-image-upscale` | [`p-image-upscale-quality-checklist.md`](https://github.com/PrunaAI/pruna-skills/tree/main/image/p-image-upscale-quality-checklist.md) | [`image-to-video`](https://github.com/PrunaAI/pruna-skills/tree/main/plugins/image-to-video/skills/image-to-video/SKILL.md), [`narrated-multi-scene`](https://github.com/PrunaAI/pruna-skills/tree/main/plugins/narrated-multi-scene/skills/narrated-multi-scene/SKILL.md), [`WORKFLOW-RECIPES`](https://github.com/PrunaAI/pruna-skills/tree/main/docs/WORKFLOW-RECIPES.md), [`generate_upscale_comparison.py`](https://github.com/PrunaAI/pruna-skills/tree/main/workflows/_shared/scripts/generate_upscale_comparison.py) |\n| `p-image-try-on` | [`p-image-try-on-quality-checklist.md`](https://github.com/PrunaAI/pruna-skills/tree/main/image/p-image-try-on-quality-checklist.md) | [`WORKFLOW-RECIPES`](https://github.com/PrunaAI/pruna-skills/tree/main/docs/WORKFLOW-RECIPES.md), [`p-image-try-on`](https://github.com/PrunaAI/pruna-skills/tree/main/plugins/p-image-try-on/skills/p-image-try-on/SKILL.md), [`realistic-persona-showcase.md`](https://github.com/PrunaAI/pruna-skills/tree/main/shared/realistic-persona-showcase.md) |\n| `p-video` | [`p-video-quality-checklist.md`](https://github.com/PrunaAI/pruna-skills/tree/main/video/p-video-quality-checklist.md) | [`image-to-video`](https://github.com/PrunaAI/pruna-skills/tree/main/plugins/image-to-video/skills/image-to-video/SKILL.md), [`narrated-multi-scene`](https://github.com/PrunaAI/pruna-skills/tree/main/plugins/narrated-multi-scene/skills/narrated-multi-scene/SKILL.md), [`interactive-explainer`](https://github.com/PrunaAI/pruna-skills/tree/main/plugins/interactive-explainer/skills/interactive-explainer/SKILL.md), [`visual-transition-reel`](https://github.com/PrunaAI/pruna-skills/tree/main/plugins/visual-transition-reel/skills/visual-transition-reel/SKILL.md), [`WORKFLOW-RECIPES`](https://github.com/PrunaAI/pruna-skills/tree/main/docs/WORKFLOW-RECIPES.md) |\n| `p-video-avatar` | [`p-video-avatar-quality-checklist.md`](https://github.com/PrunaAI/pruna-skills/tree/main/video/p-video-avatar-quality-checklist.md) · persona bar: [`realistic-persona-showcase.md`](https://github.com/PrunaAI/pruna-skills/tree/main/shared/realistic-persona-showcase.md) | [`avatar-single-scene`](https://github.com/PrunaAI/pruna-skills/tree/main/plugins/avatar-single-scene/skills/avatar-single-scene/SKILL.md), [`avatar-multi-scene`](https://github.com/PrunaAI/pruna-skills/tree/main/plugins/avatar-multi-scene/skills/avatar-multi-scene/SKILL.md), [`interactive-explainer`](https://github.com/PrunaAI/pruna-skills/tree/main/plugins/interactive-explainer/skills/interactive-explainer/SKILL.md), [`WORKFLOW-RECIPES`](https://github.com/PrunaAI/pruna-skills/tree/main/docs/WORKFLOW-RECIPES.md) |\n| `p-video-animate` | [`p-video-animate-quality-checklist.md`](https://github.com/PrunaAI/pruna-skills/tree/main/video/p-video-animate-quality-checklist.md) | [`avatar-multi-scene`](https://github.com/PrunaAI/pruna-skills/tree/main/plugins/avatar-multi-scene/skills/avatar-multi-scene/SKILL.md), [`WORKFLOW-RECIPES`](https://github.com/PrunaAI/pruna-skills/tree/main/docs/WORKFLOW-RECIPES.md) |\n| `p-video-replace` | [`p-video-replace-quality-checklist.md`](https://github.com/PrunaAI/pruna-skills/tree/main/video/p-video-replace-quality-checklist.md) | [`p-video-replace`](https://github.com/PrunaAI/pruna-skills/tree/main/plugins/p-video-replace/skills/p-video-replace/SKILL.md), [`generate_video_comparison.py`](https://github.com/PrunaAI/pruna-skills/tree/main/workflows/_shared/scripts/generate_video_comparison.py) |\n| `music-2.5` + music video assembly | [`music-video-quality-checklist.md`](https://github.com/PrunaAI/pruna-skills/tree/main/workflows/music-video-quality-checklist.md) | [`music-video`](https://github.com/PrunaAI/pruna-skills/tree/main/plugins/music-video/skills/music-video/SKILL.md), [`music-2.5`](https://github.com/PrunaAI/pruna-skills/tree/main/plugins/music-2.5/skills/music-2.5/SKILL.md) |\n\n## Core checklist (all models)\n\n- **[Generation diversity](./generation-diversity.md)** — ritual seed + rotate scenario axes on **every** model (image, video, try-on, avatar, …).\n- **[Random seed ritual](./random-seed-ritual.md) (SSoT)** — generate and state a ritual string **before** every generation; derive prompt axes via sum-mod; never copy example strings from docs.\n- Goal and acceptance criteria are explicit (what \"good\" looks like is written down).\n- Input assets are valid and licensed (URL/file reachable, rights cleared).\n- Prompt and settings match the intended output format (`aspect_ratio`, duration, resolution, style lock). **Video default:** `720p`, `24` fps unless the brief asks for final `1080p` / `48`.\n- Output contains no accidental watermarks, UI overlays, or stray text unless requested.\n- Brand, legal, and safety constraints are satisfied before handoff.\n- Manifest/log captures model, input fields, prediction id, output URL, and **`ritual_seed`** for traceability.\n\n## Model-specific checklists\n\n- [`p-image-quality-checklist.md`](https://github.com/PrunaAI/pruna-skills/tree/main/image/p-image-quality-checklist.md)\n- [`p-image-edit-quality-checklist.md`](https://github.com/PrunaAI/pruna-skills/tree/main/image/p-image-edit-quality-checklist.md)\n- [`p-image-upscale-quality-checklist.md`](https://github.com/PrunaAI/pruna-skills/tree/main/image/p-image-upscale-quality-checklist.md)\n- [`p-image-try-on-quality-checklist.md`](https://github.com/PrunaAI/pruna-skills/tree/main/image/p-image-try-on-quality-checklist.md)\n- [`p-video-quality-checklist.md`](https://github.com/PrunaAI/pruna-skills/tree/main/video/p-video-quality-checklist.md)\n- [`p-video-avatar-quality-checklist.md`](https://github.com/PrunaAI/pruna-skills/tree/main/video/p-video-avatar-quality-checklist.md)\n- [`p-video-animate-quality-checklist.md`](https://github.com/PrunaAI/pruna-skills/tree/main/video/p-video-animate-quality-checklist.md)\n- [`p-video-replace-quality-checklist.md`](https://github.com/PrunaAI/pruna-skills/tree/main/video/p-video-replace-quality-checklist.md)\n- [`music-video-quality-checklist.md`](https://github.com/PrunaAI/pruna-skills/tree/main/workflows/music-video-quality-checklist.md)\n\n## Visual variety (launch reels)\n\nBefore **any** generation, run [generation-diversity.md](./generation-diversity.md). Launch reels: also [visual-variety-bible.md](https://github.com/PrunaAI/pruna-skills/tree/main/shared/visual-variety-bible.md) **Variety checklist**. Persona/playground bar: [realistic-persona-showcase.md](https://github.com/PrunaAI/pruna-skills/tree/main/shared/realistic-persona-showcase.md).\n\nFor phased human review before expensive video jobs, see [staged-generation-gate.md](https://github.com/PrunaAI/pruna-skills/tree/main/references/policies/staged-generation-gate.md) and the per-skill index [workflow-feedback-gates.md](https://github.com/PrunaAI/pruna-skills/tree/main/references/policies/workflow-feedback-gates.md).\n\n## Workflow note\n\nFor multi-scene projects, run these checks per scene and add a final continuity pass\n(style, character identity, voice, and pacing consistency across scenes).\n\n**Narrated cinematic B-roll:** validate [scene anchor triple](https://github.com/PrunaAI/pruna-skills/tree/main/video/scene-anchor-triple.md) inputs before `p-video` — start still, end still, uploaded narration URL per row.\n\nFile v1.0.6:references/random-seed-ritual.md\n\n# Random seed ritual (SSoT — mandatory before every generation)\n\nThe random seed ritual is a lean [String Seed of Thought](https://pub.sakana.ai/ssot/) (DAG) protocol. **Every** Pruna generation — every prompt, every `POST /v1/predictions`, every scene row — starts here.\n\nThis prevents copy-pasting example strings (`k7Qm2xP9`, `482901`, …) and reduces accidental duplicate outputs across sessions.\n\n## The ritual (do this first)\n\nBefore writing prompts, curl, or runner JSON:\n\n1. **Generate a random string** in-agent (8–16 chars, mixed case + digits).\n2. **Log it** as `ritual_seed` in the manifest / internal plan. Do **not** require a user-visible *\"Ritual seed: …\"* line unless the user asks for transparency.\n3. **Derive prompt choices** from the string — sum char codes, mod N — pick axes from [generation-diversity.md](./generation-diversity.md) (`aspect_ratio`, `camera_tag`, `render_category_tag`, …).\n4. **Write the prompt** using [explicit prompt structure](./generation-diversity.md#explicit-prompt-structure-required) and derived axes.\n5. **Record** axes chosen and prediction id in the manifest alongside `ritual_seed`.\n\n**Do not pass the ritual string to API `seed`.** API runs without `seed` unless the user explicitly requests reproducibility (`api_seed`).\n\n**Never** proceed to `POST /v1/predictions` without completing steps 1–2 (unless the user supplied an explicit `api_seed` — see below).\n\n## Reuse rules\n\n| Situation | Action |\n|-----------|--------|\n| **New hero / independent still / mood-board panel** | Fresh ritual string |\n| **Same-brief slop retry** | Reuse same `ritual_seed`; note `retry_ritual_seed` in manifest |\n| **Same character arc** | Lock **hero plate URL** + cast descriptor; reuse `ritual_seed` only on same-brief regen |\n| **User says \"lock seed\" / provides integer** | Pass **their** number as `api_seed` → `input.seed`; skip new ritual for that chain |\n\nCharacter continuity = approved plate URL + cast descriptor — **not** the ritual string on the API.\n\n## Anti-patterns\n\n| Wrong | Right |\n|-------|--------|\n| Copy example strings from SKILL.md | Fresh ritual string each independent generation |\n| Pass ritual string as API `seed` | Ritual is planning-only; `api_seed` only when user asks |\n| One ritual string for entire mood board | New ritual per independent **`p-image`** |\n| Skip ritual because API `seed` is optional | Ritual always; API omits `seed` by default |\n\n## Example (internal plan / optional user-visible)\n\nManifest: `\"ritual_seed\": \"k7Qm2xP9\"`. Derived: aspect_ratio 16:9, camera_tag fish-eye, render_category_tag cartoon_anime_fantasy.  \nPrompt: Disco ball reflections on an otter DJ scratching vinyl at a packed 1970s roller rink, fish-eye lens, glitter confetti mid-air, funky energy.  \n…then curl / runner **without** `\"seed\"` in `input`.\n\n## Manifest snippet\n\n```json\n{\n  \"ritual_seed_policy\": \"ssot_dag_before_every_generation\",\n  \"ritual_seed\": \"k7Qm2xP9\",\n  \"seed_log\": [\n    { \"phase\": \"hero_p_image\", \"ritual_seed\": \"k7Qm2xP9\", \"creature_tag\": \"otter_dj\", \"setting_tag\": \"1970s_roller_rink\", \"prompt_hash\": \"…\" },\n    { \"phase\": \"scene_2_avatar\", \"ritual_seed\": \"k7Qm2xP9\", \"scene_id\": 2 }\n  ]\n}\n```\n\n## Where this applies\n\nAll Pruna generation skills and workflow runners — **every invocation**:\n\n- **`p-image`**, **`p-image-edit`**, **`p-image-try-on`**, **`p-image-upscale`**\n- **`p-video`**, **`p-video-avatar`**, **`p-video-animate`**, **`p-video-replace`**\n\n## Related\n\n- [generation-diversity.md](./generation-diversity.md) — ritual + axis rotation + sum-mod derivation\n- [realistic-persona-showcase.md](https://github.com/PrunaAI/pruna-skills/tree/main/shared/realistic-persona-showcase.md) — persona planning\n- [staged-generation-gate.md](https://github.com/PrunaAI/pruna-skills/tree/main/references/policies/staged-generation-gate.md) — approval phases\n- [approval-red-flags.md](https://github.com/PrunaAI/pruna-skills/tree/main/references/policies/approval-red-flags.md) — red flags\n\nFile v1.0.6:references/replicate-api.md\n\n# Replicate API (minimal)\n\nUsed by external tool skills (e.g. [stable-audio-2.5](../SKILL.md), [gemini-3.1-flash-tts](https://github.com/PrunaAI/pruna-skills/tree/main/plugins/gemini-3.1-flash-tts/skills/gemini-3.1-flash-tts/SKILL.md)).\n\n**Missing token:** agents must stop and point the user to [api-credentials.md](./api-credentials.md) — sign up at [replicate.com/account/api-tokens](https://replicate.com/account/api-tokens) ([sign in](https://replicate.com/signin) if needed), then `export REPLICATE_API_TOKEN=r8_...`.\n\n## Auth\n\n```bash\nexport REPLICATE_API_TOKEN=r8_...\n```\n\nHeader: `Authorization: Bearer ${REPLICATE_API_TOKEN}`\n\n## Create + poll\n\n```bash\n# POST https://api.replicate.com/v1/models/{owner}/{name}/predictions\n# Body: {\"input\": { ... }}\n\n# Poll GET on response.urls.get until status == succeeded\n# Download output URL (string or list depending on model)\n```\n\nShared client: [`workflows/_shared/scripts/replicate_api.py`](https://github.com/PrunaAI/pruna-skills/tree/main/workflows/_shared/scripts/replicate_api.py)\n\n## Stable Audio 2.5\n\nModel: `stability-ai/stable-audio-2.5`  \nRequired input: `prompt`  \nOptional: `duration` (1–190), `steps` (4–8), `cfg_scale`, `seed`\n\n## Music 2.5 (MiniMax)\n\nModel: `minimax/music-2.5`  \nRequired input: `lyrics` (1–3,500 chars, structure tags supported)  \nOptional: `prompt` (style), `sample_rate`, `bitrate`, `audio_format` (`mp3` default)\n\nWorkflow: [music-video](https://github.com/PrunaAI/pruna-skills/tree/main/plugins/music-video/skills/music-video/SKILL.md) · tool skill: [music-2.5](https://github.com/PrunaAI/pruna-skills/tree/main/plugins/music-2.5/skills/music-2.5/SKILL.md)\n\n## Gemini 3.1 Flash TTS\n\nModel: `google/gemini-3.1-flash-tts`  \nRequired input: `text`  \nOptional: `voice` (default `Kore`), `prompt` (style/scene), `language_code` (default `en-US`)\n\nOutput: audio file URL. Use for narration — upload to Pruna as part of [scene anchor triple](https://github.com/PrunaAI/pruna-skills/tree/main/workflows/scene-anchor-triple/SKILL.md) (`input.audio` + `input.image` + `input.last_frame_image` on `p-video`). Layering with beds: [audio-post-production.md](./audio-post-production.md)\n\nFile v1.0.6:README-INSTALL.md\n\n# stable-audio-2.5\n\n## Install\n\n**Skills CLI** (copy-paste):\n\n```bash\nnpx skills add PrunaAI/pruna-skills@stable-audio-2.5 -y\n```\n\n**Plugins CLI** (workflow bundles with deps — pick from the list):\n\n```bash\nnpx plugins add PrunaAI/pruna-skills\n# when prompted, select: stable-audio-2.5\n```\n\nDo **not** use `npx plugins add PrunaAI/pruna-skills@stable-audio-2.5` — the plugins CLI has no `@name` filter (that’s skills only) and prints “No plugins found”.\n\n**Claude Code:**\n\n```text\n/plugin marketplace add PrunaAI/pruna-skills\n/plugin install stable-audio-2.5@pruna-skills\n```\n\nList all skills:\n\n```bash\nnpx skills add PrunaAI/pruna-skills -l\n```\n\nAfter install, start a **new chat**. See the [root README](https://github.com/PrunaAI/pruna-skills/tree/main/workflows/README/SKILL.md).\n\n## From a local clone\n\n```bash\nnpx skills add .@stable-audio-2.5 -y\n# or:\nnpx skills add ./plugins/stable-audio-2.5/skills --skill stable-audio-2.5 -y\n```\n\nFile v1.0.6:skill-card.md\n\n## Description: <br>\nUse when someone wants light instrumental background music: an ambient bed under dialogue or underscore for reels and explainers. <br>\n\nThis skill is ready for commercial/non-commercial use. <br>\n\n## Publisher: <br>\n[pruna-ai](https://clawhub.ai/user/pruna-ai) <br>\n\n### License/Terms of Use: <br>\nMIT-0 <br>\n\n\n## Use Case: <br>\nDevelopers and content creators use this skill to generate light instrumental MP3 background beds with Replicate Stable Audio 2.5 and mix them under launch reels, explainers, or narrated video. <br>\n\n### Deployment Geography for Use: <br>\nGlobal <br>\n\n## Known Risks and Mitigations: <br>\nRisk: The package includes broad image and video generation guidance outside the stated background-music purpose. <br>\nMitigation: Review before installing and constrain use to Stable Audio background-bed generation unless the broader workflow guidance has been separately approved. <br>\nRisk: Replicate predictions are paid API calls and require credential handling. <br>\nMitigation: Check that REPLICATE_API_TOKEN is present before calling the API, stop when it is missing, and avoid printing or embedding secrets in prompts, logs, manifests, or task text. <br>\nRisk: Generated music may interfere with narration or dialogue if prompt and mix settings are too aggressive. <br>\nMitigation: Use instrumental, no-vocal prompts, keep bed volume low, and inspect the generated audio or final mix before advancing. <br>\n\n\n## Reference(s): <br>\n- [Stable Audio 2.5 on Replicate](https://replicate.com/stability-ai/stable-audio-2.5) <br>\n- [Replicate API](references/replicate-api.md) <br>\n- [API Credentials](references/api-credentials.md) <br>\n- [Audio Post-Production](references/audio-post-production.md) <br>\n- [Random Seed Ritual](references/random-seed-ritual.md) <br>\n- [Generation Diversity](references/generation-diversity.md) <br>\n- [Generation Quality Checklists](references/generation-quality-checklists.md) <br>\n\n\n## Skill Output: <br>\n**Output Type(s):** [text, markdown, code, shell commands, configuration, guidance] <br>\n**Output Format:** [Markdown guidance with JSON and shell command snippets] <br>\n**Output Parameters:** [1D] <br>\n**Other Properties Related to Output:** [Guides an agent to call Replicate for an MP3 URL and optionally mix downloaded audio under video with ffmpeg.] <br>\n\n## Skill Version(s): <br>\n1.0.6 (source: server release metadata and skill frontmatter metadata) <br>\n\n## Ethical Considerations: <br>\nUsers should evaluate whether this skill is appropriate for their environment, review any generated or modified files before relying on them, and apply their organization's safety, security, and compliance requirements before deployment. <br>\n\nFile v1.0.6:skill.manifest.json\n\n{\n  \"scripts\": {\n    \"core\": [],\n    \"shared\": []\n  },\n  \"references\": [\n    \"replicate-api.md\",\n    \"api-credentials.md\",\n    \"audio-post-production.md\"\n  ]\n}\n\nArchive v1.0.2: 8 files, 10406 bytes\n\nFiles: README-INSTALL.md (291b), references/api-credentials.md (3138b), references/audio-post-production.md (9313b), references/replicate-api.md (2178b), skill-card.md (2526b), skill.manifest.json (160b), SKILL.md (4034b), _meta.json (135b)\n\nFile v1.0.2:SKILL.md\n\n---\nname: stable-audio-2.5\ndescription: Use when the user wants light instrumental background music, an ambient bed under dialogue or voiceover, or underscore for reels and explainers.\nlicense: MIT\nmetadata:\n  version: \"1.0.2\"\n  provider: replicate\n  replicate_model: stability-ai/stable-audio-2.5\n---\n\n# Stable Audio 2.5 (Replicate)\n\nText-to-music for **instrumental background beds** on launch reels. Not a Pruna P-model — runs on [Replicate](https://replicate.com/stability-ai/stable-audio-2.5).\n\n**Mix helper (repo):** [`launch_background_music.py`](https://github.com/PrunaAI/pruna-skills/tree/main/workflows/_shared/scripts/launch_background_music.py) — probes video length, generates bed, mixes under VO with ffmpeg.\n\n## When to use\n\n| Goal | Use this |\n|------|----------|\n| Light instrumental under a concat launch reel | Yes — after final assembly |\n| Under **embedded narration** from scene anchor triple | Yes — mix quiet bed after concat — [audio-post-production.md](./references/audio-post-production.md) |\n| Replace avatar VO | No — bed mixes **under** existing dialogue |\n| Pruna-native audio | No — use [`p-video`](https://github.com/PrunaAI/pruna-skills/tree/main/plugins/p-video/skills/p-video/SKILL.md) audio input instead |\n\n## Environment\n\n```bash\nexport REPLICATE_API_TOKEN=r8_...\n```\n\nRequires **`ffmpeg`** and **`ffprobe`** on PATH for mix step.\n\n## Model input (Replicate)\n\n| Field | Notes |\n|-------|-------|\n| `prompt` | **Required.** Style tags work well — e.g. *Instrumental light electronic pop, soft groove, mellow synth pads, no vocals, 94 BPM* |\n| `duration` | Seconds, 1–190 (match or slightly exceed reel length) |\n| `steps` | 4–8 (default 8) |\n| `cfg_scale` | 1–25 (default 1) |\n| `seed` | Optional integer for reproducibility |\n\nOutput: single **MP3** URL.\n\n## HTTP (curl)\n\n### Create prediction\n\n```bash\ncurl -s -X POST \\\n  -H \"Authorization: Bearer ${REPLICATE_API_TOKEN}\" \\\n  -H \"Content-Type: application/json\" \\\n  -d '{\n    \"input\": {\n      \"prompt\": \"Instrumental light electronic pop bed, soft groove and mellow synth pads, calm positive tech atmosphere, understated background music, no vocals, 94 BPM\",\n      \"duration\": 90,\n      \"steps\": 8,\n      \"cfg_scale\": 1\n    }\n  }' \\\n  \"https://api.replicate.com/v1/models/stability-ai/stable-audio-2.5/predictions\"\n```\n\nPoll `urls.get` until `status` is `succeeded`; download `output` MP3.\n\n## Launch reel integration\n\n### Plan JSON (`background_music`)\n\n```json\n\"background_music\": {\n  \"enabled\": true,\n  \"prompt\": \"Instrumental light electronic pop bed, soft groove and mellow synth pads, calm positive tech atmosphere, understated background music, no vocals, 94 BPM\",\n  \"volume\": 0.12,\n  \"output_name\": \"skills_library_announcement_with_music.mp4\"\n}\n```\n\n### Post-concat bed on any MP4\n\n```bash\npython3 workflows/_shared/scripts/launch_background_music.py \\\n  --video output/my_reel/final.mp4 \\\n  --volume 0.12\n```\n\n## Prompt tips (launch beds)\n\n- Lead with **Instrumental** and **no vocals**\n- Name mood: light, calm, positive, soft groove, mellow — keep *understated* and *background music* so beds sit under VO\n- Optional BPM (**88–98** for tech launch reels; default script uses **94**)\n- Avoid *energetic*, *driving*, or very high BPM — beds should support dialogue, not compete with it\n- Avoid lyrics, song title, or artist name triggers\n- Scene-specific beds (restaurant, jungle, boutique): name the setting but keep groove soft and BPM in the high 80s–mid 90s\n\n## Related\n\n- [audio-post-production.md](./references/audio-post-production.md) — narration + bed layering\n- [gemini-3.1-flash-tts](https://github.com/PrunaAI/pruna-skills/tree/main/plugins/gemini-3.1-flash-tts/skills/gemini-3.1-flash-tts/SKILL.md) — narration voiceover\n- [visual-transition-reel](https://github.com/PrunaAI/pruna-skills/tree/main/plugins/visual-transition-reel/skills/visual-transition-reel/SKILL.md) — concat + optional bed assembly\n- [replicate-api.md](./references/replicate-api.md) — shared Replicate patterns\n\nFile v1.0.2:_meta.json\n\n{\n  \"ownerId\": \"kn7cagwf7q3t0cxrgteb7xk0bh81j0eb\",\n  \"slug\": \"stable-audio-2-5\",\n  \"version\": \"1.0.2\",\n  \"publishedAt\": 1784208916613\n}\n\nFile v1.0.2:references/api-credentials.md\n\n# API credentials (Pruna + Replicate)\n\n**Agent rule:** Before any `POST /v1/predictions`, Replicate prediction, or paid runner — check env vars. If a required key is **missing or empty**, **stop** and tell the user how to sign up. Do not guess, mock, or skip with placeholder keys.\n\n## Pruna P-API\n\n| | |\n|--|--|\n| **Env var** | `PRUNA_API_KEY` |\n| **Header** | `apikey: ${PRUNA_API_KEY}` (not `Authorization: Bearer`) |\n| **Sign up / get key** | [Pruna dashboard](https://dashboard.pruna.ai/) |\n| **Docs** | [Quickstart](https://docs.api.pruna.ai/guides/quickstart) · [pruna-api.md](https://github.com/PrunaAI/pruna-skills/tree/main/workflows/core/pruna-api/SKILL.md) |\n\n**Used by:** all `p-image*`, `p-video*` tool skills and Pruna workflow runners.\n\n### If `PRUNA_API_KEY` is missing — agent message template\n\n> Pruna generation needs an API key. Sign up or sign in at **[dashboard.pruna.ai](https://dashboard.pruna.ai/)**, create an API key, then set:\n>\n> ```bash\n> export PRUNA_API_KEY=\"your_key_here\"\n> ```\n>\n> Add that to your shell profile or project `.env` (never commit the key). Reply when it’s set and we can continue.\n\n## Replicate\n\n| | |\n|--|--|\n| **Env var** | `REPLICATE_API_TOKEN` |\n| **Header** | `Authorization: Bearer ${REPLICATE_API_TOKEN}` |\n| **Sign up / get token** | [Replicate API tokens](https://replicate.com/account/api-tokens) ([sign in](https://replicate.com/signin) first if needed) |\n| **Docs** | [replicate-api.md](./replicate-api.md) |\n\n**Used by:** `music-2.5`, `gemini-3.1-flash-tts`, `stable-audio-2.5`, `whisperx`, and workflow beds/TTS/song phases.\n\n### If `REPLICATE_API_TOKEN` is missing — agent message template\n\n> This step uses Replicate (song, TTS, transcription, or background bed). Create a token at **[replicate.com/account/api-tokens](https://replicate.com/account/api-tokens)**, then set:\n>\n> ```bash\n> export REPLICATE_API_TOKEN=\"r8_...\"\n> ```\n>\n> Reply when it’s set and we can continue.\n\n## Which key does this job need?\n\n| Task | Keys required |\n|------|----------------|\n| `p-image`, `p-image-edit`, `p-image-upscale`, `p-image-try-on` | `PRUNA_API_KEY` |\n| `p-video`, `p-video-avatar`, `p-video-animate`, `p-video-replace` | `PRUNA_API_KEY` |\n| Music 2.5 song generation | `REPLICATE_API_TOKEN` |\n| Gemini TTS narration | `REPLICATE_API_TOKEN` |\n| Stable Audio background bed | `REPLICATE_API_TOKEN` |\n| WhisperX transcription | `REPLICATE_API_TOKEN` |\n| Music video / explainer (full pipeline) | **Both** — Pruna for stills/video; Replicate for song/TTS/bed as needed |\n\nWhen only one key is missing, suggest **only** that provider’s signup link — not both.\n\n## Security\n\n- Never print full keys in chat or commit them to git.\n- `.env` is gitignored; prefer env vars over hardcoding in plans or manifests.\n- Never embed keys in prompts, manifests, plan JSON, logs, or **subagent task text**.\n- Prefer the **parent agent** to own API calls; do not fan credentials across parallel subagents unless the host documents isolated secret injection.\n- Full rules: [agent-safety.md](https://github.com/PrunaAI/pruna-skills/tree/main/workflows/core/agent-safety/SKILL.md).\n\nFile v1.0.2:references/audio-post-production.md\n\n# Audio post-production (Pruna + Replicate)\n\nHow to choose and **layer** audio when building reels, multi-scene films, and launch videos.\n\n**Multi-scene narrated films:** use the [scene anchor triple](https://github.com/PrunaAI/pruna-skills/tree/main/workflows/core/scene-anchor-triple/SKILL.md) — pass TTS to **`p-video`** as `input.audio` with `image` + `last_frame_image`; do not post-mux unless re-render is impossible.\n\n**Visual-only transitions (no VO):** use the [scene anchor pair](https://github.com/PrunaAI/pruna-skills/tree/main/workflows/core/scene-anchor-pair/SKILL.md) — `duration` instead of `audio`; see [visual-transition-reel](https://github.com/PrunaAI/pruna-skills/tree/main/plugins/visual-transition-reel/skills/visual-transition-reel/SKILL.md).\n\n## Audio-led `p-video` (required when VO/narration exists)\n\nWhen narration, TTS, or a timed audio slice is available **before** video render:\n\n1. Upload the audio file to Pruna (`POST /v1/files`).\n2. Pass `urls.get` as **`input.audio`** on **`p-video`** (or **`p-video-avatar`** for human lip-sync).\n3. **Omit `duration`** — clip length follows the audio (capped at **20s** on P-API); the model syncs motion to speech.\n4. Set **`save_audio`: true** so the full line is embedded in the output clip.\n5. **Probe TTS length** before render — per-scene lines should be **≤ ~19s** or the API truncates the tail even when `audio` is set.\n5. **Concat** clips in order (narration already on each clip). Optional bed mixed **under** VO in post.\n\n**Never** generate silent `p-video` and ffmpeg-mux narration afterward unless re-render is impossible — post-mux **truncates** lines longer than the video slot (common with Gemini TTS).\n\n**Over 20s?** Shorten scene copy → tighten TTS pace in `style_prompt` → split into two scene rows (each with its own triple). See [narrated-multi-scene](https://github.com/PrunaAI/pruna-skills/tree/main/plugins/narrated-multi-scene/skills/narrated-multi-scene/SKILL.md) duration gate.\n\nHelper: [`p_video_payload.py`](https://github.com/PrunaAI/pruna-skills/tree/main/workflows/_shared/scripts/p_video_payload.py) — `build_p_video_payload(...)` enforces omitting `duration` when `audio_url` is set.\n\n| Workflow | Audio source | `p-video` fields |\n|----------|--------------|------------------|\n| Dog plush / story film | Gemini TTS per scene | `image` + `last_frame_image` + `audio` |\n| Music video performance | Song slice per cut | `image` + `audio` |\n| Music video B-roll | Song slice (optional) | `image` + `audio` or `duration` only |\n| Viking narrator beats | Gemini TTS | `image` + `last_frame_image` + `audio` |\n\n## Tool picker\n\n| Need | Tool | Skill |\n|------|------|-------|\n| Cinematic clip with model-generated sound | `p-video` (`save_audio`, optional uploaded `audio`) | [p-video](https://github.com/PrunaAI/pruna-skills/tree/main/plugins/p-video/skills/p-video/SKILL.md) |\n| Lip-sync / duration locked to VO | Upload audio → `p-video` with `audio` | [p-video](https://github.com/PrunaAI/pruna-skills/tree/main/plugins/p-video/skills/p-video/SKILL.md) |\n| Documentary / story narrator | [Gemini 3.1 Flash TTS](https://replicate.com/google/gemini-3.1-flash-tts) | [gemini-3.1-flash-tts](https://github.com/PrunaAI/pruna-skills/tree/main/plugins/gemini-3.1-flash-tts/skills/gemini-3.1-flash-tts/SKILL.md) |\n| Light instrumental under dialogue | [Stable Audio 2.5](https://replicate.com/stability-ai/stable-audio-2.5) | [stable-audio-2.5](../SKILL.md) |\n| Full song with sung vocals | [Music 2.5](https://replicate.com/minimax/music-2.5) | [music-2.5](https://github.com/PrunaAI/pruna-skills/tree/main/plugins/music-2.5/skills/music-2.5/SKILL.md) |\n| Speaking on-camera character | `p-video-avatar` | [p-video-avatar](https://github.com/PrunaAI/pruna-skills/tree/main/plugins/p-video-avatar/skills/p-video-avatar/SKILL.md) |\n\n**Env:** Pruna calls need `PRUNA_API_KEY`; Replicate audio tools need `REPLICATE_API_TOKEN`. Assembly steps need **`ffmpeg`** / **`ffprobe`**.\n\n## Layering matrix\n\n| Stack | Primary audio | Secondary | Mix notes |\n|-------|---------------|-----------|-----------|\n| **Silent B-roll** | — | — | Concat video only |\n| **Native `p-video` sound** | Model output | — | Keep `save_audio` default; normalize in assembly if scenes differ |\n| **Narration only (fallback)** | Gemini TTS | — | Post-mux only when audio-led `p-video` is not suitable — prefer **Pipeline B** below |\n| **Bed only** | Stable Audio bed | — | [`launch_background_music.py`](https://github.com/PrunaAI/pruna-skills/tree/main/workflows/_shared/scripts/launch_background_music.py) |\n| **Narration + bed (preferred)** | Gemini TTS → **`p-video` `audio`** | Stable Audio (quiet) | TTS uploaded to Pruna drives clip length + sync; bed mixed in post under narration (~0.08–0.15) |\n| **Avatar VO + bed** | `p-video-avatar` dialogue | Stable Audio bed | Same bed pattern as replace/launch reels — bed **under** existing speech |\n| **Music video** | Music 2.5 full song | — | [music-video](https://github.com/PrunaAI/pruna-skills/tree/main/plugins/music-video/skills/music-video/SKILL.md) |\n\n## Recommended pipelines\n\n### A — Narrated multi-scene B-roll (**preferred — scene anchor triple**)\n\n```text\nPhase 0 — intake: scene table with start/end still prompts + narration lines\nPhase 1 — hero + p-image-edit start stills + end stills (parallel)\nPhase 2 — Gemini TTS per scene (parallel) → upload each to /v1/files\nPhase 3 — p-video per scene: input.image + input.last_frame_image + input.audio (parallel; omit duration)\nPhase 4 — ffmpeg concat (VO embedded; frame chain via shared end/start URLs)\nPhase 5 — optional Stable Audio bed under narration\n```\n\n**Scene anchor triple:** same pattern as first/last frame pairing — `audio` is the third required upload per scene row. **`p-video-avatar`:** portrait + optional `last_frame_image` + uploaded `audio`.\n\n### A′ — Post-mux narration (fallback only)\n\nUse only when you already have silent clips and cannot re-render. Risk: TTS longer than clip slots → cut-off VO.\n\n```text\nPhase 3 — p-video I2V without audio → concat → mux TTS in ffmpeg\n```\n\n### C — Launch / product reel (existing pattern)\n\n```text\nPhase 1 — p-video-avatar or replace reel → concat\nPhase 2 — Stable Audio bed via launch_background_music.py (bed under VO, not replacing it)\n```\n\n## ffmpeg mixing (conceptual)\n\n**Narration onto silent concat** (single VO file):\n\n```bash\nffmpeg -y -i concat_video.mp4 -i narration.mp3 \\\n  -map 0:v -map 1:a -c:v copy -c:a aac -b:a 192k -shortest output_with_vo.mp4\n```\n\n**Bed under existing narration + video** (same pattern as `launch_background_music.py`):\n\n```text\n[1:a]volume=0.12,aloop=...[bed];\n[0:a][bed]amix=inputs=2:duration=first[aout]\n```\n\nNarration / avatar dialogue stays on stream `0:a`; bed is stream `1:a` at low volume.\n\n**Bed on silent concat** — loop a short generated clip to full video length (no per-assemble Stable Audio call):\n\n```text\n[1:a]volume=0.12,aloop=loop=-1:size=2e+09[bed]  →  map video + [bed], -shortest\n```\n\nPlan field `\"reuse_bed\": true` skips regeneration when `audio/launch_bed.mp3` exists. Delete that file (or set `reuse_bed: false`) only when you want a new prompt or seed.\n\n## Intake questions (audio)\n\nAsk before generating paid audio or video:\n\n| Topic | Questions |\n|-------|-----------|\n| **Primary voice** | Narrator (Gemini TTS), on-screen avatar (`p-video-avatar`), or native `p-video` sound only? |\n| **Narration scope** | Per-scene lines vs one continuous VO track? |\n| **Music / bed** | None, instrumental bed only, or full song ([music-2.5](https://github.com/PrunaAI/pruna-skills/tree/main/plugins/music-2.5/skills/music-2.5/SKILL.md))? |\n| **Sync strategy** | **Preferred:** TTS → Pruna upload → **`p-video` / `p-video-avatar` with `audio`** (clip length = audio). Post-mux only as fallback. |\n| **Levels** | Bed volume target (default ~0.12 under avatar VO; ~0.08–0.12 under Gemini narration)? |\n\n## Manifest fields\n\n```json\n{\n  \"narration\": { \"enabled\": true, \"voice\": \"Sulafat\", \"mode\": \"per_scene\" },\n  \"background_music\": { \"enabled\": true, \"reuse_bed\": true, \"volume\": 0.10, \"prompt\": \"Instrumental ... no vocals\" },\n  \"p_video_audio\": { \"save_audio\": true }\n}\n```\n\n## Limitations (from [P-Video on Replicate](https://replicate.com/prunaai/p-video))\n\n- Native SFX/dialogue quality varies — for premium voice realism, prefer **Gemini TTS** or **`p-video-avatar`**, then optionally mix a bed.\n- Multi-speaker native audio can drift; dedicated TTS per role is safer for narration-heavy cuts.\n- Extreme camera motion and complex multi-scene stories are weaker than **frame-anchored chaining** + per-scene prompts — see [p-video](https://github.com/PrunaAI/pruna-skills/tree/main/plugins/p-video/skills/p-video/SKILL.md) **First / last frame chaining**.\n\n## Related\n\n- [parallel-execution.md](https://github.com/PrunaAI/pruna-skills/tree/main/workflows/core/parallel-execution/SKILL.md) — phased vs parallel when frames chain\n- [narrated-multi-scene](https://github.com/PrunaAI/pruna-skills/tree/main/plugins/narrated-multi-scene/skills/narrated-multi-scene/SKILL.md)\n- [pruna-generative-pipeline](https://github.com/PrunaAI/pruna-skills/tree/main/plugins/pruna-generative-pipeline/skills/pruna-generative-pipeline/SKILL.md)\n\nFile v1.0.2:references/replicate-api.md\n\n# Replicate API (minimal)\n\nUsed by external tool skills (e.g. [stable-audio-2.5](../SKILL.md), [gemini-3.1-flash-tts](https://github.com/PrunaAI/pruna-skills/tree/main/plugins/gemini-3.1-flash-tts/skills/gemini-3.1-flash-tts/SKILL.md)).\n\n**Missing token:** agents must stop and point the user to [api-credentials.md](./api-credentials.md) — sign up at [replicate.com/account/api-tokens](https://replicate.com/account/api-tokens) ([sign in](https://replicate.com/signin) if needed), then `export REPLICATE_API_TOKEN=r8_...`.\n\n## Auth\n\n```bash\nexport REPLICATE_API_TOKEN=r8_...\n```\n\nHeader: `Authorization: Bearer ${REPLICATE_API_TOKEN}`\n\n## Create + poll\n\n```bash\n# POST https://api.replicate.com/v1/models/{owner}/{name}/predictions\n# Body: {\"input\": { ... }}\n\n# Poll GET on response.urls.get until status == succeeded\n# Download output URL (string or list depending on model)\n```\n\nShared client: [`workflows/_shared/scripts/replicate_api.py`](https://github.com/PrunaAI/pruna-skills/tree/main/workflows/_shared/scripts/replicate_api.py)\n\n## Stable Audio 2.5\n\nModel: `stability-ai/stable-audio-2.5`  \nRequired input: `prompt`  \nOptional: `duration` (1–190), `steps` (4–8), `cfg_scale`, `seed`\n\n## Music 2.5 (MiniMax)\n\nModel: `minimax/music-2.5`  \nRequired input: `lyrics` (1–3,500 chars, structure tags supported)  \nOptional: `prompt` (style), `sample_rate`, `bitrate`, `audio_format` (`mp3` default)\n\nWorkflow: [music-video](https://github.com/PrunaAI/pruna-skills/tree/main/plugins/music-video/skills/music-video/SKILL.md) · tool skill: [music-2.5](https://github.com/PrunaAI/pruna-skills/tree/main/plugins/music-2.5/skills/music-2.5/SKILL.md)\n\n## Gemini 3.1 Flash TTS\n\nModel: `google/gemini-3.1-flash-tts`  \nRequired input: `text`  \nOptional: `voice` (default `Kore`), `prompt` (style/scene), `language_code` (default `en-US`)\n\nOutput: audio file URL. Use for narration — upload to Pruna as part of [scene anchor triple](https://github.com/PrunaAI/pruna-skills/tree/main/workflows/core/scene-anchor-triple/SKILL.md) (`input.audio` + `input.image` + `input.last_frame_image` on `p-video`). Layering with beds: [audio-post-production.md](./audio-post-production.md)\n\nFile v1.0.2:README-INSTALL.md\n\n# stable-audio-2.5\n\nCopy `SKILL.md` into your agent skills folder or install via the repo plugin manifest.\n\nRequires `REPLICATE_API_TOKEN`, `ffmpeg`, and `ffprobe`.\n\nMix script ships with workflow bundles that list `launch_background_music.py` in `skill.manifest.json` → `scripts.shared`.\n\nFile v1.0.2:skill-card.md\n\n## Description: <br>\nUse when the user wants light instrumental background music, an ambient bed under dialogue or voiceover, or underscore for reels and explainers. <br>\n\nThis skill is ready for commercial/non-commercial use. <br>\n\n## Publisher: <br>\n[pruna-ai](https://clawhub.ai/user/pruna-ai) <br>\n\n### License/Terms of Use: <br>\nMIT-0 <br>\n\n\n## Use Case: <br>\nDevelopers and creators use this skill to generate light instrumental MP3 background beds with Replicate Stable Audio 2.5 and mix them quietly under launch reels, explainers, dialogue, or voiceover. <br>\n\n### Deployment Geography for Use: <br>\nGlobal <br>\n\n## Known Risks and Mitigations: <br>\nRisk: Prompts and generated audio requests are sent to Replicate for hosted processing. <br>\nMitigation: Use the skill only for content approved for Replicate processing and avoid private client details, secrets, or confidential unreleased content in prompts. <br>\nRisk: Replicate API tokens could be exposed if placed in files, prompts, logs, or generated plans. <br>\nMitigation: Keep REPLICATE_API_TOKEN in environment variables, do not print full tokens, and stop when the token is missing rather than using placeholders. <br>\nRisk: The mix step depends on local ffmpeg and ffprobe availability. <br>\nMitigation: Confirm ffmpeg and ffprobe are on PATH before attempting post-production mixing. <br>\n\n\n## Reference(s): <br>\n- [Replicate Stable Audio 2.5 model](https://replicate.com/stability-ai/stable-audio-2.5) <br>\n- [replicate-api.md](artifact/references/replicate-api.md) <br>\n- [api-credentials.md](artifact/references/api-credentials.md) <br>\n- [audio-post-production.md](artifact/references/audio-post-production.md) <br>\n- [ClawHub skill page](https://clawhub.ai/pruna-ai/skills/stable-audio-2-5) <br>\n\n\n## Skill Output: <br>\n**Output Type(s):** [guidance, markdown, shell commands, configuration, API calls] <br>\n**Output Format:** [Markdown guidance with inline JSON and bash examples] <br>\n**Output Parameters:** [1D] <br>\n**Other Properties Related to Output:** [Produces Replicate prediction requests and an MP3 output URL; optional post-processing uses ffmpeg and ffprobe.] <br>\n\n## Skill Version(s): <br>\n1.0.2 (source: server release metadata and SKILL.md metadata) <br>\n\n## Ethical Considerations: <br>\nUsers should evaluate whether this skill is appropriate for their environment, review any generated or modified files before relying on them, and apply their organization's safety, security, and compliance requirements before deployment. <br>\n\nFile v1.0.2:skill.manifest.json\n\n{\n  \"scripts\": {\n    \"core\": [],\n    \"shared\": []\n  },\n  \"references\": [\n    \"replicate-api.md\",\n    \"api-credentials.md\",\n    \"audio-post-production.md\"\n  ]\n}","readmeExcerpt":"Skill: stable-audio-2.5 Owner: pruna-ai Summary: Use when someone wants light instrumental background music — an ambient bed under dialogue or underscore for reels and explainers. Tags: ai:1.0.14, generative:1.0.14, latest:1.0.14, pruna:1.0.14 Version history: v1.0.14 | 2026-09-29T15:36:14.516Z | auto - Updated version metadata from 1.0.13 to 1.0.14 in SKILL.md. - Removed the file skill-card.md. v1.0.13 | 2026-09-17T","codeSnippets":[],"executableExamples":[{"language":"bash","snippet":"export REPLICATE_API_TOKEN=r8_..."},{"language":"bash","snippet":"curl -s -X POST \\\n  -H \"Authorization: Bearer ${REPLICATE_API_TOKEN}\" \\\n  -H \"Content-Type: application/json\" \\\n  -d '{"},{"language":"bash","snippet":"curl -s -X POST \\\n  -H \"Authorization: Bearer ${REPLICATE_API_TOKEN}\" \\\n  -H \"Content-Type: application/json\" \\\n  -d '{\n    \"input\": {\n      \"prompt\": \"Instrumental light electronic pop bed, soft groove and mellow synth pads, calm positive tech atmosphere, understated background music, no vocals, 94 BPM\",\n      \"duration\": 90,\n      \"steps\": 8,\n      \"cfg_scale\": 1\n    }\n  }' \\\n  \"https://api.replicate.com/v1/models/stability-ai/stable-audio-2.5/predictions\""},{"language":"bash","snippet":"export REPLICATE_API_TOKEN=r8_..."},{"language":"bash","snippet":"curl -s -X POST \\\n  -H \"Authorization: Bearer ${REPLICATE_API_TOKEN}\" \\\n  -H \"Content-Type: application/json\" \\\n  -d '{"},{"language":"bash","snippet":"curl -s -X POST \\\n  -H \"Authorization: Bearer ${REPLICATE_API_TOKEN}\" \\\n  -H \"Content-Type: application/json\" \\\n  -d '{\n    \"input\": {\n      \"prompt\": \"Instrumental light electronic pop bed, soft groove and mellow synth pads, calm positive tech atmosphere, understated background music, no vocals, 94 BPM\",\n      \"duration\": 90,\n      \"steps\": 8,\n      \"cfg_scale\": 1\n    }\n  }' \\\n  \"https://api.replicate.com/v1/models/stability-ai/stable-audio-2.5/predictions\""}],"parameters":null,"dependencies":[],"permissions":[],"extractedFiles":[{"path":"SKILL.md","content":"---\nname: stable-audio-2.5\ndescription: Use when someone wants light instrumental background music — an ambient bed under dialogue or underscore for reels and explainers.\nlicense: MIT\nmetadata:\n  version: \"1.0.14\"\n  package: pruna-skills\n  provider: replicate\n  replicate_model: stability-ai/stable-audio-2.5\n---\n\n## Prerequisites\n\nInstall and load these skills before generating (skip if already in context via `@pruna`):\n\n| Skill | Description | Install |\n| --- | --- | --- |\n| `generation-diversity` | Use when writing any generative prompt — ritual seed, explicit structure, scenario axes, and quality gates before paid API calls. | `npx skills add PrunaAI/pruna-skills@generation-diversity -y` |\n| `audio-prompting` | Use when crafting TTS, music, or bed prompts for any generative audio model — director style, song structure, and post-production layering. | `npx skills add PrunaAI/pruna-skills@audio-prompting -y` |\n| `pruna-api` | Use before any Pruna or Replicate HTTP call — credentials, upload/poll/download, parallel batches, and agent safety. | `npx skills add PrunaAI/pruna-skills@pruna-api -y` |\n\nOr install the full suite once: `npx skills add PrunaAI/pruna-skills@pruna -y`\n\nFollow each skill's **Before generating** / craft sections — do not restate guide content here.\n\n## Agent habit\n\nIn the **first reply**, name `` `stable-audio-2.5` `` in backticks, confirm `REPLICATE_API_TOKEN` (or stop with signup links from `pruna-api`), then ask for required inputs. Open intake → **`generation-diversity`** clarification intake before the first `POST`. Redirect when **When NOT to use** fits better.\n\n## When NOT to use\n\nUse a different skill instead:\n\n| Skill | Description | Install |\n| --- | --- | --- |\n| `music-2.5` | Use when someone wants an original AI song with vocals — sung lyrics, a style prompt track, or source audio for a music video. | `npx skills add PrunaAI/pruna-skills@music-2.5 -y` |\n| `gemini-3.1-flash-tts` | Use when someone needs spoken narration or voiceover — explainer tracks, documentary lines, or voice to pair with generated video. | `npx skills add PrunaAI/pruna-skills@gemini-3.1-flash-tts -y` |\n\n## Environment\n\n```bash\nexport REPLICATE_API_TOKEN=r8_...\n```\n\nRequires **`ffmpeg`** and **`ffprobe`** on PATH for the mix step.\n\n## HTTP (curl)\n\n```bash\ncurl -s -X POST \\\n  -H \"Authorization: Bearer ${REPLICATE_API_TOKEN}\" \\\n  -H \"Content-Type: application/json\" \\\n  -d '{\n    \"input\": {\n      \"prompt\": \"Instrumental light electronic pop bed, soft groove and mellow synth pads, calm positive tech atmosphere, understated background music, no vocals, 94 BPM\",\n      \"duration\": 90,\n      \"steps\": 8,\n      \"cfg_scale\": 1\n    }\n  }' \\\n  \"https://api.replicate.com/v1/models/stability-ai/stable-audio-2.5/predictions\"\n```\n\nPoll `urls.get` until `status` is `succeeded`; download `output` MP3.\n\n## Before generating\n\n1. Complete Prerequisites guide reading order — bed prompt craft: `audio-prompting` **Worked examples** (instrumental bed).\n2. Confirm **`pro"},{"path":"_meta.json","content":"{\n  \"ownerId\": \"kn7cagwf7q3t0cxrgteb7xk0bh81j0eb\",\n  \"slug\": \"stable-audio-2-5\",\n  \"version\": \"1.0.14\",\n  \"publishedAt\": 1790696174516\n}"},{"path":"skill-card.md","content":"## Description:\n\nHelps create light instrumental background music for dialogue, reels, and explainers using Stable Audio 2.5.\n\nThis skill is ready for commercial/non-commercial use.\n\n## Publisher:\n\n[pruna-ai](https://clawhub.ai/user/pruna-ai)\n\n### License/Terms of Use:\n\nMIT-0\n\n## Use Case:\n\nCreators and editors use this skill to prepare instrumental music prompts and generation steps for unobtrusive beds beneath narration or short videos.\n\n### Deployment Geography for Use:\n\nGlobal\n\n## Known Risks and Mitigations:\n\nRisk: Unpinned skill installation commands may fetch changed dependencies.\n\nMitigation: Review the installation commands and use a trusted or pinned installation path when available.\n\nRisk: Prompts or credentials may be exposed when calling an external audio service.\n\nMitigation: Provide REPLICATE_API_TOKEN only when using Replicate, and omit secrets and private information from prompts.\n\n## Reference(s):\n\n- [ClawHub skill release](https://clawhub.ai/pruna-ai/skills/stable-audio-2-5)\n\n## Skill Output:\n\n**Output Type(s):** [Text, Shell commands, Configuration guidance]\n\n**Output Format:** [Markdown with bash and JSON examples]\n\n**Output Parameters:** [1D]\n\n**Other Properties Related to Output:** [Guides Replicate MP3 generation; optional mixing requires ffmpeg and ffprobe.]\n\n## Skill Version(s):\n\n1.0.14 (source: skill frontmatter and server release)\n\n## Ethical Considerations:\n\nUsers should evaluate whether this skill is appropriate for their environment, review any generated or modified files before relying on them, and apply their organization's safety, security, and compliance requirements before deployment."},{"path":"skill.manifest.json","content":"{\n  \"references\": []\n}"}],"languages":[],"docsSourceLabel":"CLAWHUB","editorialOverview":"Use when someone wants light instrumental background music — an ambient bed under dialogue or underscore for reels and explainers. Skill: stable-audio-2.5 Owner: pruna-ai Summary: Use when someone wants light instrumental background music — an ambient bed under dialogue or underscore for reels and explainers. Tags: ai:1.0.14, generative:1.0.14, latest:1.0.14, pruna:1.0.14 Version history: v1.0.14 | 2026-09-29T15:36:14.516Z | auto - Updated version metadata from 1.0.13 to 1.0.14 in SKILL.md. - Removed the file skill-card.md. v1.0.13 | 2026-09-17T","editorialQuality":{"score":100,"threshold":65,"status":"ready","wordCount":1191,"uniquenessScore":47,"reasons":[]}},"media":{"evidence":{"source":"no-media","verified":false,"confidence":"low","updatedAt":"2026-10-10T12:32:11.870Z","emptyReason":"No screenshots, media assets, or demo links are available."},"primaryImageUrl":null,"mediaAssetCount":0,"assets":[],"demoUrl":null},"ownerResources":{"evidence":{"source":"unclaimed","verified":false,"confidence":"low","updatedAt":"2026-10-10T12:32:11.870Z","emptyReason":"This page has not been claimed by the agent owner."},"hasCustomPage":false,"customPageUpdatedAt":null,"customLinks":[],"structuredLinks":{"docsUrl":null,"demoUrl":null,"supportUrl":null,"pricingUrl":null,"statusUrl":null},"customPage":null},"relatedAgents":{"evidence":{"source":"protocol-neighbors","verified":false,"confidence":"medium","updatedAt":"2026-10-10T14:46:32.564Z","emptyReason":null},"items":[{"id":"8ebccd8e-3863-4187-8355-c3f14e1f9edf","entityType":"agent","canonicalPath":"/agent/iofficeai-aionui","slug":"iofficeai-aionui","name":"AionUi","description":"Free, local, open-source 24/7 Cowork app and OpenClaw for Gemini CLI, Claude Code, Codex, OpenCode, Qwen Code, Goose CLI, Auggie, and more | 🌟 Star if you like it!","url":"https://github.com/iOfficeAI/AionUi","homepage":"https://www.aionui.com","source":"GITHUB_REPOS","protocols":["MCP","OPENCLAW"],"capabilities":[],"safetyScore":100,"overallRank":70,"updatedAt":"2026-10-09T19:11:12.944Z","createdAt":"2026-02-25T03:38:16.584Z","downloads":null},{"id":"b917f68a-ebff-438e-84f8-3f4b2494c0bc","entityType":"agent","canonicalPath":"/agent/activepieces-activepieces","slug":"activepieces-activepieces","name":"activepieces","description":"AI Agents & MCPs & AI Workflow Automation • (~400 MCP servers for AI agents) • AI Automation / AI Agent with MCPs • AI Workflows & AI Agents • MCPs for AI Agents","url":"https://github.com/activepieces/activepieces","homepage":"https://www.activepieces.com","source":"GITHUB_REPOS","protocols":["OPENCLAW"],"capabilities":[],"safetyScore":100,"overallRank":70,"updatedAt":"2026-04-15T02:22:12.426Z","createdAt":"2026-02-25T03:38:12.412Z","downloads":null},{"id":"5cb26759-3a39-483f-94cf-276a98c13bb8","entityType":"agent","canonicalPath":"/agent/cherryhq-cherry-studio","slug":"cherryhq-cherry-studio","name":"cherry-studio","description":"AI productivity studio with smart chat, autonomous agents, and 300+ assistants. Unified access to frontier LLMs","url":"https://github.com/CherryHQ/cherry-studio","homepage":"https://cherry-ai.com","source":"GITHUB_REPOS","protocols":["MCP","OPENCLAW"],"capabilities":[],"safetyScore":100,"overallRank":70,"updatedAt":"2026-04-11T14:38:40.986Z","createdAt":"2026-02-25T03:38:19.379Z","downloads":null},{"id":"6f6582d0-5d76-4f0f-b81d-86520247950b","entityType":"agent","canonicalPath":"/agent/copilotkit-copilotkit","slug":"copilotkit-copilotkit","name":"CopilotKit","description":"The Frontend for Agents & Generative UI. React + Angular","url":"https://github.com/CopilotKit/CopilotKit","homepage":"https://docs.copilotkit.ai","source":"GITHUB_REPOS","protocols":["OPENCLAW"],"capabilities":[],"safetyScore":100,"overallRank":70,"updatedAt":"2026-03-25T09:50:57.846Z","createdAt":"2026-02-25T03:39:14.617Z","downloads":null}],"links":{"hub":"/agent","source":"/agent/source/clawhub","protocols":[{"label":"OpenClaw","href":"/agent/protocol/openclew"}]}}}