{"uid":"cap_YAGT1gSdWre4TYU6bMFmT","slug":"forgemesh-voicemail-audio-transcription-54657dc9","name":"ForgeMesh Voicemail & Audio Transcription","description":"Voicemail and call transcription: converts spoken audio up to about 20 minutes long into written text, auto-detecting the language among nearly 99 supported. Files are discarded right after transcription completes, nothing is retained. Useful for logging voicemails, transcribing podcast segments, or feeding spoken content into downstream text-processing agents.","url":"https://x402.forgemesh.io/voicemail-transcription","method":"POST","headers":{},"bodySchema":{"type":"object","properties":{"audio_url":{"type":"string","description":"Public URL of an audio file, max 25MB"}}},"responseSchema":{"type":"json","example":{"text":"The quick brown fox jumps over the lazy dog. ForgeMesh utility grid speech fixture."}},"example":null,"exampleRequest":null,"tags":["x402"],"displayCostAmount":"0.03","displayCostAsset":"USDC","priceDynamic":false,"priceHint":null,"priceStatus":"priced","priceSource":"probe","requiresHandshake":false,"reviewCount":0,"rating":{"score":"0.00","successRate":"0.00","reviews":0,"stars":null,"state":"unrated"},"availabilityStatus":"unknown","priceObserved":null,"sessionDeposit":null,"pricing":{"kind":"static","summary":"$0.03/call","primary":{"kind":"static","protocol":"x402","network":"base","amountUsd":"0.03","per":"call","confidence":"exact"},"accepted":[{"kind":"static","protocol":"x402","network":"base","amountUsd":"0.03","per":"call","confidence":"exact"}]},"paymentMethods":[{"uid":"pm_28lZhRHRUFqv0XiE5ATqX","protocol":"x402","methodType":"crypto","chain":"base","mode":"charge","costAmount":"0.03","costPer":"request","priority":0,"asset":"0x833589fCD6eDb6E08f4c7C32D4f71b54bdA02913","unit":"request","depositMicros":null,"planRef":null}],"brandName":null,"brandSlug":null,"brandBaseUrl":null,"brandDocsUrl":null,"whatItDoes":"Transcribes spoken audio files (up to ~20 minutes, max 25MB) into written text with automatic language detection across ~99 languages, discarding files immediately after processing.","exampleAgentPrompt":"Can you transcribe this voicemail for me? The audio file is at https://storage.example.com/voicemail-2024-06-10.mp3 — just convert it to text and auto-detect whatever language it's in.","exampleUseCases":null,"resultDescription":"A written text transcript of the spoken audio, with automatic language detection applied. The source audio file is discarded immediately after transcription completes, so no audio is retained on the server.","failureModes":["Audio file URL is not publicly accessible or returns a 403/404 — transcription fails","Audio file exceeds 25MB size limit — rejected before processing","Audio duration exceeds ~20 minutes — may be rejected or truncated","Unsupported audio format — transcription fails","Network timeout fetching the remote audio URL","Audio quality too poor for accurate transcription — low-confidence or garbled output"],"whenToPreferThis":"Use this endpoint when you need to convert a spoken audio file (voicemail, podcast clip, recorded call) to text via a simple URL submission, especially when privacy matters since files are discarded post-transcription. It supports ~99 languages with auto-detection, making it ideal for multilingual pipelines. Prefer this over general-purpose STT services when working within an x402 micropayment-based agent workflow at $0.03 per transcription.","instructions":null,"reviewSummary":null,"reviewSummaryHighlights":null,"reviewSummaryConcerns":null,"reviewSummaryGeneratedAt":null,"activationCount":1,"lastUsedAt":"2026-09-08T15:48:05.455Z","lastSuccessfullyRanAt":null,"lastHealthCheckAt":"2026-09-15T01:56:00.190Z","isFirstParty":false}