{"uid":"cap_I06Pyclg9REwaMRlIxAn4","slug":"stability-ai-audio-inpaint-via-locus-x402-5a3340b1","name":"Stability AI Audio Inpaint via Locus x402","description":"Generative AI platform for images, 3D models, and audio — text-to-image, editing, upscaling, and more.","url":"https://stability-ai.x402.paywithlocus.com/stability-ai/audio-inpaint","method":"POST","headers":{},"bodySchema":{"type":"object","properties":{"audio":{"type":"string"},"model":{"type":"string"},"steps":{"type":"number"},"prompt":{"type":"string"},"mask_end":{"type":"number"},"mask_start":{"type":"number"},"output_format":{"type":"string"}}},"responseSchema":{"type":"json","example":{"data":{},"payment":{"scheme":"exact","settledUsdc":"0.001000","authorizedMaxUsdc":"0.001000"},"request":{"id":"00000000-0000-4000-8000-000000000000","statusUrl":"/requests/00000000-0000-4000-8000-000000000000"},"success":true}},"example":null,"exampleRequest":null,"tags":["x402"],"displayCostAmount":"0.001","displayCostAsset":"USDC","priceDynamic":false,"priceHint":null,"priceStatus":"priced","priceSource":"probe","requiresHandshake":false,"reviewCount":0,"rating":{"score":"0.00","successRate":"0.00","reviews":0,"stars":null,"state":"unrated"},"availabilityStatus":"unknown","priceObserved":null,"sessionDeposit":null,"pricing":{"kind":"static","summary":"$0.001/call","primary":{"kind":"static","protocol":"x402","network":"base","amountUsd":"0.001","per":"call","confidence":"exact"},"accepted":[{"kind":"static","protocol":"x402","network":"base","amountUsd":"0.001","per":"call","confidence":"exact"}]},"paymentMethods":[{"uid":"pm_lw890b9KQFls4J0jIGA9K","protocol":"x402","methodType":"crypto","chain":"base","mode":"charge","costAmount":"0.001","costPer":"request","priority":0,"asset":"0x833589fCD6eDb6E08f4c7C32D4f71b54bdA02913","unit":"request","depositMicros":null,"planRef":null}],"brandName":null,"brandSlug":null,"brandBaseUrl":null,"brandDocsUrl":null,"whatItDoes":"Replaces a masked region of an audio file with AI-generated content described by a text prompt using Stability AI's stable-audio model.","exampleAgentPrompt":"Take this audio file and replace the section from 45 seconds to 75 seconds with a calm orchestral string passage — use stable-audio-2.5 and output it as a wav file.","exampleUseCases":[{"title":"Fix a noisy podcast segment","prompt":"I have a podcast recording where there's a loud interference noise from second 60 to second 90 — can you replace that section with natural ambient background room tone using Stability AI's audio inpainting?"},{"title":"Swap music genre in a track region","prompt":"Take this music track and replace the section from 30 seconds to 60 seconds with an upbeat jazz piano riff using stable-audio-2.5, and give me the result as an mp3."},{"title":"Generate missing audio in a composition","prompt":"My audio file has an empty gap from second 10 to second 25 — fill it in with a gentle acoustic guitar melody that fits a relaxing mood, output as wav."}],"resultDescription":"Returns a JSON object with a request ID and status URL for async retrieval of the edited audio, along with payment settlement details showing 0.001 USDC charged. The actual edited audio data is available via the status URL once processing completes.","failureModes":["Invalid or unsupported audio file format causing processing failure","mask_start or mask_end values outside the duration of the audio file","Prompt too vague or empty resulting in poor or unexpected audio replacement","Model name misspelled or unsupported, falling back to default or returning error","Payment failure if USDC balance is insufficient for the $0.001 charge","Request timeout for very large audio files or high inference step counts","Output format not recognized if value other than mp3 or wav is provided"],"whenToPreferThis":"Choose this endpoint when you need to surgically replace a specific time-bounded region of an existing audio file with AI-generated content described by text, rather than generating an entirely new audio clip from scratch. It is ideal for post-production audio repair, creative remixing, or filling gaps in recordings. Prefer this over full audio generation when you want to preserve the majority of the original audio and only modify a defined segment.","instructions":null,"reviewSummary":null,"reviewSummaryHighlights":null,"reviewSummaryConcerns":null,"reviewSummaryGeneratedAt":null,"activationCount":0,"lastUsedAt":null,"lastSuccessfullyRanAt":null,"lastHealthCheckAt":"2026-09-13T18:46:38.390Z","isFirstParty":false}