{"uid":"cap_PwVNsx89iGcYeZnQkZESb","slug":"jarvisclaw-audio-transcription-whisper-large-v3-37614cdf","name":"JarvisClaw Audio Transcription (Whisper Large v3)","description":"Transcribes audio to text with whisper-large-v3. Server-side fetches the audio URL (max 25 MB), relays it to Venice's audio/transcriptions endpoint, and returns the transcript…","url":"https://api.jarvisclaw.ai/v1/marketplace/api/audio-transcribe","method":"GET","headers":{},"bodySchema":{"type":"object","properties":{"query":{"type":"string","description":"Request parameters"}}},"responseSchema":null,"example":null,"exampleRequest":null,"tags":["x402"],"displayCostAmount":"0.0115","displayCostAsset":"USDC","priceDynamic":false,"priceHint":null,"priceStatus":"priced","priceSource":"probe","requiresHandshake":false,"reviewCount":0,"rating":{"score":"0.00","successRate":"0.00","reviews":0,"stars":null,"state":"unrated"},"availabilityStatus":"unknown","priceObserved":null,"sessionDeposit":null,"pricing":{"kind":"static","summary":"$0.0115/call","primary":{"kind":"static","protocol":"x402","network":"base","amountUsd":"0.0115","per":"call","confidence":"exact"},"accepted":[{"kind":"static","protocol":"x402","network":"base","amountUsd":"0.0115","per":"call","confidence":"exact"}]},"paymentMethods":[{"uid":"pm_fycykUo2kq0tfmQddEJVd","protocol":"x402","methodType":"crypto","chain":"base","mode":"charge","costAmount":"0.0115","costPer":"request","priority":0,"asset":"0x833589fCD6eDb6E08f4c7C32D4f71b54bdA02913","unit":"request","depositMicros":null,"planRef":null}],"brandName":null,"brandSlug":null,"brandBaseUrl":null,"brandDocsUrl":null,"whatItDoes":"Transcribes audio files to text using OpenAI's Whisper large-v3 model, server-side fetching the audio from a URL and returning the transcript via Venice's transcription backend.","exampleAgentPrompt":"Can you transcribe this audio recording for me? The file is at https://example.com/interview.mp3 — it's about 10 minutes long and I need the full text transcript.","exampleUseCases":[{"title":"Podcast episode to text","prompt":"I have a podcast episode hosted at https://mypodcast.com/ep42.mp3 — can you transcribe the whole thing into text so I can turn it into a blog post?"},{"title":"Meeting recording transcript","prompt":"Transcribe this Zoom recording I uploaded to my storage bucket at https://storage.example.com/meeting-2024-06.mp3 so I can pull out the action items."},{"title":"Voicemail to text conversion","prompt":"I've got a voicemail file at https://myserver.com/voicemail.wav — can you convert it to text so I can read what was said?"}],"resultDescription":"A plain-text transcript of the spoken content in the audio file, converted using Whisper large-v3 via Venice AI's transcriptions endpoint. The response contains the transcribed text extracted from the provided audio URL.","failureModes":["Audio file exceeds 25 MB limit — request rejected","Invalid or inaccessible audio URL — server cannot fetch the file","Unsupported audio format — transcription fails","Network error fetching the remote audio URL","Venice backend unavailable or rate-limited — upstream failure","Payment not provided or insufficient USDC — 402 response returned"],"whenToPreferThis":"Choose this endpoint when you need pay-per-call audio transcription without managing API keys for OpenAI or Whisper directly, especially in agentic/x402 workflows where micropayments in USDC on Base are preferred. Ideal when the audio is already accessible via a public URL and you want server-side fetching handled for you. Best for files under 25 MB where you need Whisper large-v3 quality.","instructions":null,"reviewSummary":null,"reviewSummaryHighlights":null,"reviewSummaryConcerns":null,"reviewSummaryGeneratedAt":null,"activationCount":0,"lastUsedAt":null,"lastSuccessfullyRanAt":null,"lastHealthCheckAt":"2026-09-14T06:30:11.193Z","isFirstParty":false}