{"uid":"cap_EmOnPogy8fhP4qnR2RUG1","slug":"minichan-ocr-nvidia-nim-minimax-m3-image-text-extraction-22c84e78","name":"MiniChan OCR — NVIDIA NIM MiniMax-M3 Image Text Extraction","description":"NVIDIA NIM MiniMax-M3 (428B MoE) powered API for AI agents. 15 multimodal endpoints — text, vision, code, reasoning, creative. Pay with USDC on Base via x402.","url":"https://minizzzan.vercel.app/api/ocr","method":"POST","headers":{},"bodySchema":{"type":"object","required":["image_url"],"properties":{"image_url":{"type":"string","description":"URL of the image to analyze"}}},"responseSchema":{"type":"object"},"example":null,"exampleRequest":null,"tags":["x402"],"displayCostAmount":"0.15","displayCostAsset":"USDC","priceDynamic":false,"priceHint":null,"priceStatus":"priced","priceSource":"probe","requiresHandshake":false,"reviewCount":0,"rating":{"score":"0.00","successRate":"0.00","reviews":0,"stars":null,"state":"unrated"},"availabilityStatus":"unknown","priceObserved":null,"sessionDeposit":null,"pricing":{"kind":"static","summary":"$0.15/call","primary":{"kind":"static","protocol":"x402","network":"base","amountUsd":"0.15","per":"call","confidence":"exact"},"accepted":[{"kind":"static","protocol":"x402","network":"base","amountUsd":"0.15","per":"call","confidence":"exact"}]},"paymentMethods":[{"uid":"pm_HXnr6SUOI2FKHpIeQOWjk","protocol":"x402","methodType":"crypto","chain":"base","mode":"charge","costAmount":"0.15","costPer":"request","priority":0,"asset":"0x833589fCD6eDb6E08f4c7C32D4f71b54bdA02913","unit":"request","depositMicros":null,"planRef":null}],"brandName":null,"brandSlug":null,"brandBaseUrl":null,"brandDocsUrl":null,"whatItDoes":"Extracts and recognizes text from an image URL using NVIDIA NIM's MiniMax-M3 multimodal model","exampleAgentPrompt":"Can you read and extract all the text from this image for me? Here's the URL: https://example.com/receipt.jpg","exampleUseCases":[{"title":"Extract receipt data for expense reports","prompt":"I need to digitize this receipt image for our expense tracking system. Can you read all the text from this photo and pull out the vendor name, date, and total amount? Here's the URL: https://example.com/receipt.jpg"},{"title":"Transcribe handwritten notes from photos","prompt":"I took a photo of my handwritten notes from a meeting but need them in text form. Can you extract all the handwriting from this image so I can search and edit it? The image is at: https://example.com/notes.jpg"},{"title":"Read text from multilingual documents","prompt":"This is a scanned document with text in multiple languages. Can you extract and recognize all the text from this image, including any non-English content? Here's the URL: https://example.com/multilingual_doc.jpg"}],"resultDescription":"Returns a JSON object containing the text extracted from the provided image, powered by NVIDIA NIM's MiniMax-M3 multimodal model (428B MoE). Includes recognized text content parsed from the visual content of the image.","failureModes":["Invalid or inaccessible image URL returns an error","Image with no readable text may return empty or minimal results","Very low resolution or heavily distorted images may yield inaccurate OCR","Unsupported image formats may cause processing failures","Network timeout if the image URL is slow to load","Payment failure (insufficient USDC balance) blocks the request"],"whenToPreferThis":"Choose this endpoint when you need to extract text from images using a state-of-the-art large multimodal model (MiniMax-M3 428B MoE via NVIDIA NIM). It is particularly well-suited for complex documents, mixed-language text, or images where smaller OCR models struggle. Prefer this over lightweight OCR APIs when accuracy on difficult images is critical and you can pay per-call in USDC on Base.","instructions":null,"reviewSummary":null,"reviewSummaryHighlights":null,"reviewSummaryConcerns":null,"reviewSummaryGeneratedAt":null,"activationCount":0,"lastUsedAt":null,"lastSuccessfullyRanAt":null,"lastHealthCheckAt":"2026-09-15T18:36:56.294Z","isFirstParty":false}