{"uid":"cap_YPdeSWE1EfvonFKZGc_P7","slug":"minichan-image-analysis-nvidia-nim-minimax-m3-df040801","name":"MiniChan Image Analysis — NVIDIA NIM MiniMax-M3","description":"NVIDIA NIM MiniMax-M3 (428B MoE) powered API for AI agents. 15 multimodal endpoints — text, vision, code, reasoning, creative. Pay with USDC on Base via x402.","url":"https://app-minichwaan.vercel.app/api/analyze-image","method":"POST","headers":{},"bodySchema":{"type":"object","required":["image_url"],"properties":{"prompt":{"type":"string","description":"What to analyze in the image (default: full description)"},"image_url":{"type":"string","description":"URL of the image to analyze"}}},"responseSchema":{"type":"object"},"example":null,"exampleRequest":null,"tags":["x402"],"displayCostAmount":"0.1","displayCostAsset":"USDC","priceDynamic":false,"priceHint":null,"priceStatus":"priced","priceSource":"registry","requiresHandshake":false,"reviewCount":0,"rating":{"score":"0.00","successRate":"0.00","reviews":0,"stars":null,"state":"unrated"},"availabilityStatus":"unknown","priceObserved":null,"sessionDeposit":null,"pricing":{"kind":"static","summary":"$0.1/call","primary":{"kind":"static","protocol":"x402","network":"base","amountUsd":"0.1","per":"call","confidence":"exact"},"accepted":[{"kind":"static","protocol":"x402","network":"base","amountUsd":"0.1","per":"call","confidence":"exact"}]},"paymentMethods":[{"uid":"pm_3XbrdZ-sCOA0lfdzxFg2Z","protocol":"x402","methodType":"crypto","chain":"base","mode":"charge","costAmount":"0.1","costPer":"request","priority":0,"asset":null,"unit":"request","depositMicros":null,"planRef":null}],"brandName":null,"brandSlug":null,"brandBaseUrl":null,"brandDocsUrl":null,"whatItDoes":"Analyzes an image from a URL using NVIDIA NIM's MiniMax-M3 428B multimodal model and returns a detailed description or targeted analysis based on an optional prompt.","exampleAgentPrompt":"Can you analyze this image for me — https://example.com/photo.jpg — and tell me what objects, people, and scene are visible in it?","exampleUseCases":[{"title":"Product image content audit","prompt":"I have a product image at https://store.example.com/item123.jpg — can you describe everything visible in it, including the product, background, and any text or labels shown?"},{"title":"Accessibility alt-text generation","prompt":"Generate a detailed description of this image so I can use it as alt-text for a visually impaired user: https://myblog.com/uploads/hero-banner.png"},{"title":"Security and compliance image check","prompt":"Look at this uploaded screenshot at https://cdn.example.com/screenshot.png and tell me if it contains any personally identifiable information like names, addresses, or ID numbers."}],"resultDescription":"Returns a JSON object containing the model's textual analysis of the image — typically a detailed description of visible content, objects, people, text, or scene context, optionally focused on the provided prompt.","failureModes":["Invalid or inaccessible image URL returns an error","Image URL pointing to a non-image resource may fail or produce poor results","Very large images may time out or be rejected","Ambiguous or missing prompt may yield a generic description rather than targeted analysis","Payment failure or insufficient USDC balance will block the request"],"whenToPreferThis":"Choose this endpoint when you need a powerful 428B-parameter multimodal model to analyze image content from a URL with flexible, prompt-guided analysis. It is well-suited for agents that require nuanced scene understanding, object detection, text extraction from images, or accessibility descriptions, and that can pay per-call in USDC on Base via x402.","instructions":null,"reviewSummary":null,"reviewSummaryHighlights":null,"reviewSummaryConcerns":null,"reviewSummaryGeneratedAt":null,"activationCount":0,"lastUsedAt":null,"lastSuccessfullyRanAt":null,"lastHealthCheckAt":"2026-09-13T18:43:25.479Z","isFirstParty":false}