{"uid":"cap_KaQxQLo9Y32kTmFQ0UpNK","slug":"forgemesh-audio-transcription-api-f0eae5fa","name":"ForgeMesh Audio Transcription API","description":"Audio transcription API for agents: paid per call, no keys, no accounts. Neural speech-recognition transcript of any public audio URL up to 25MB, with segment text. Nothing retained.","url":"https://x402.forgemesh.io/audio-transcription-api","method":"POST","headers":{},"bodySchema":{"type":"object","properties":{"audio_url":{"type":"string","description":"Public URL of an audio file, max 25MB"}}},"responseSchema":{"type":"json","example":{"text":"The quick brown fox jumps over the lazy dog. ForgeMesh utility grid speech fixture."}},"example":null,"exampleRequest":null,"tags":["x402"],"displayCostAmount":"0.03","displayCostAsset":"USDC","priceDynamic":false,"priceHint":null,"priceStatus":"priced","priceSource":"probe","requiresHandshake":false,"reviewCount":0,"rating":{"score":"0.00","successRate":"0.00","reviews":0,"stars":null,"state":"unrated"},"availabilityStatus":"unknown","priceObserved":null,"sessionDeposit":null,"pricing":{"kind":"static","summary":"$0.03/call","primary":{"kind":"static","protocol":"x402","network":"base","amountUsd":"0.03","per":"call","confidence":"exact"},"accepted":[{"kind":"static","protocol":"x402","network":"base","amountUsd":"0.03","per":"call","confidence":"exact"}]},"paymentMethods":[{"uid":"pm_WAmkeM5_HLwA9TTlexBCB","protocol":"x402","methodType":"crypto","chain":"base","mode":"charge","costAmount":"0.03","costPer":"request","priority":0,"asset":"0x833589fCD6eDb6E08f4c7C32D4f71b54bdA02913","unit":"request","depositMicros":null,"planRef":null}],"brandName":null,"brandSlug":null,"brandBaseUrl":null,"brandDocsUrl":null,"whatItDoes":"Transcribes any public audio URL (up to 25MB) into text with segment-level detail using neural speech recognition, paid per call via x402 with no account required.","exampleAgentPrompt":"Can you transcribe this audio file for me? Here's the public URL: https://example.com/interview-recording.mp3 — I need the full transcript with segments.","exampleUseCases":null,"resultDescription":"A neural speech-recognition transcript of the submitted audio file, broken into segment-level text chunks. No data is retained after processing. Response includes the full transcribed text and individual segment texts.","failureModes":["Audio file exceeds 25MB limit — request rejected","Audio URL is not publicly accessible or returns a non-200 response — fetch failure","Unsupported audio format — transcription error","Payment of $0.03 USDC not completed via x402 — 402 Payment Required response","URL is malformed or missing — validation error","Audio content is too noisy or low quality — degraded or empty transcript"],"whenToPreferThis":"Choose this endpoint when you need a fast, one-off audio transcription without setting up API keys or accounts — ideal for agents that occasionally need speech-to-text and want a pay-per-call model via x402. Best for public audio URLs under 25MB where no data retention is acceptable or desired.","instructions":null,"reviewSummary":null,"reviewSummaryHighlights":null,"reviewSummaryConcerns":null,"reviewSummaryGeneratedAt":null,"activationCount":2,"lastUsedAt":"2026-07-25T16:52:04.216Z","lastSuccessfullyRanAt":null,"lastHealthCheckAt":"2026-09-15T01:08:58.027Z","isFirstParty":false}