{"uid":"cap_HdXj3jW36_0AZ714cIIlr","slug":"whisper-edge-audio-transcription-e7518cba","name":"Whisper Edge Audio Transcription","description":"Transcribe a short audio clip (<=60s) using Workers AI Whisper. Returns verbatim text + segment timestamps. DO NOT use for copyright audio. — $0.005 USDC/call on Base. Pay-per-call via x402 — no API key, no account. Free samples daily.","url":"https://defi-intel-agent-gateway.gg-neo15.workers.dev/v1/audio/whisper-edge-transcript","method":"POST","headers":{},"bodySchema":null,"responseSchema":null,"example":null,"exampleRequest":null,"tags":["x402"],"displayCostAmount":"0.005","displayCostAsset":"USDC","priceDynamic":false,"priceHint":null,"priceStatus":"priced","priceSource":"probe","requiresHandshake":false,"reviewCount":0,"rating":{"score":"0.00","successRate":"0.00","reviews":0,"stars":null,"state":"unrated"},"availabilityStatus":"unknown","priceObserved":null,"sessionDeposit":null,"pricing":{"kind":"static","summary":"$0.005/call","primary":{"kind":"static","protocol":"x402","network":"base","amountUsd":"0.005","per":"call","confidence":"exact"},"accepted":[{"kind":"static","protocol":"x402","network":"base","amountUsd":"0.005","per":"call","confidence":"exact"}]},"paymentMethods":[{"uid":"pm_Xj3h6HozdUaPC1VtINRWM","protocol":"x402","methodType":"crypto","chain":"base","mode":"charge","costAmount":"0.005","costPer":"request","priority":0,"asset":"0x833589fCD6eDb6E08f4c7C32D4f71b54bdA02913","unit":"request","depositMicros":null,"planRef":null}],"brandName":null,"brandSlug":null,"brandBaseUrl":null,"brandDocsUrl":null,"whatItDoes":"Transcribes short audio clips (up to 60 seconds) using Workers AI Whisper, returning verbatim text and segment-level timestamps.","exampleAgentPrompt":"Can you transcribe this audio clip for me? It's a 45-second voice recording of a meeting snippet — I need the full verbatim text plus timestamps for each spoken segment.","exampleUseCases":[{"title":"Voice memo to meeting notes","prompt":"I recorded a 50-second voice memo after my call and need it transcribed into text with timestamps so I can copy it into my notes — can you run it through Whisper?"},{"title":"Podcast excerpt captioning","prompt":"I have a 30-second clip from a podcast episode that I want to caption — please transcribe it and give me the segment timestamps so I know exactly when each phrase was spoken."},{"title":"Multilingual interview snippet","prompt":"Can you transcribe this 55-second audio clip from an interview? I need the verbatim text and timestamps for each segment so I can align it with my video editor."}],"resultDescription":"Returns a JSON object containing the full verbatim transcript text of the audio clip and an array of segment-level timestamps indicating when each portion of speech occurred within the recording.","failureModes":["Audio clip exceeds 60-second limit — request rejected or truncated","Unsupported audio format — processing fails with an error","Poor audio quality or heavy background noise — low-accuracy or garbled transcript","Payment not fulfilled via x402 — call blocked before processing","Copyright-protected audio submitted — usage policy violation (not enforced technically but flagged in terms)","Network or edge worker timeout — no response returned"],"whenToPreferThis":"Choose this endpoint when you need fast, pay-per-call audio transcription without account setup, especially for short clips (<=60s) where per-call micropayment via x402 on Base is acceptable. Prefer it over full-scale ASR services when you want no API key overhead and need segment timestamps alongside verbatim text. Avoid if your clip is longer than 60 seconds or involves copyright-protected material.","instructions":null,"reviewSummary":null,"reviewSummaryHighlights":null,"reviewSummaryConcerns":null,"reviewSummaryGeneratedAt":null,"activationCount":0,"lastUsedAt":null,"lastSuccessfullyRanAt":null,"lastHealthCheckAt":"2026-09-14T12:51:05.820Z","isFirstParty":false}