{"uid":"cap_rWlV4ivwpakzELNJ8vMX_","slug":"minichan-code-from-image-nvidia-nim-minimax-m3-beb2fca3","name":"MiniChan Code-from-Image (NVIDIA NIM MiniMax-M3)","description":"NVIDIA NIM MiniMax-M3 (428B MoE) powered API for AI agents. 15 multimodal endpoints — text, vision, code, reasoning, creative. Pay with USDC on Base via x402.","url":"https://minichan.vercel.app/api/code-from-image","method":"POST","headers":{},"bodySchema":{"type":"object","required":["image_url"],"properties":{"image_url":{"type":"string","description":"URL of the image to analyze"}}},"responseSchema":{"type":"object"},"example":null,"exampleRequest":null,"tags":["x402"],"displayCostAmount":"0.2","displayCostAsset":"USDC","priceDynamic":false,"priceHint":null,"priceStatus":"priced","priceSource":"unknown","requiresHandshake":false,"reviewCount":0,"rating":{"score":"0.00","successRate":"0.00","reviews":0,"stars":null,"state":"unrated"},"availabilityStatus":"unknown","priceObserved":null,"sessionDeposit":null,"pricing":{"kind":"static","summary":"$0.2/call","primary":{"kind":"static","protocol":"x402","network":"base","amountUsd":"0.2","per":"call","confidence":"exact"},"accepted":[{"kind":"static","protocol":"x402","network":"base","amountUsd":"0.2","per":"call","confidence":"exact"}]},"paymentMethods":[{"uid":"pm_bfXVYcylb9suq-BS1AIkB","protocol":"x402","methodType":"crypto","chain":"base","mode":"charge","costAmount":"0.2","costPer":"request","priority":0,"asset":null,"unit":"request","depositMicros":null,"planRef":null}],"brandName":null,"brandSlug":null,"brandBaseUrl":null,"brandDocsUrl":null,"whatItDoes":"Analyzes an image URL and generates corresponding source code using NVIDIA NIM's MiniMax-M3 428B MoE multimodal model","exampleAgentPrompt":"Look at this screenshot of a login form — https://example.com/login-mockup.png — and generate the HTML, CSS, and JavaScript code to build it.","exampleUseCases":null,"resultDescription":"A JSON object containing generated source code derived from analyzing the provided image, likely including code snippets, language identification, and implementation details based on the visual content of the image.","failureModes":["Invalid or inaccessible image URL returns an error","Image content not interpretable as code-related visual returns generic or empty code","Payment failure via x402/USDC on Base blocks the request","Model timeout on complex images with large MoE inference","Unsupported image format or corrupted image URL"],"whenToPreferThis":"Choose this endpoint when you need to convert visual representations — screenshots, wireframes, mockups, diagrams, or whiteboard photos — directly into source code using a powerful 428B parameter multimodal model. Prefer this over generic OCR or simpler vision models when you need actual runnable code output rather than text extraction. Ideal for developers automating UI-to-code workflows or reverse-engineering interfaces from images.","instructions":null,"reviewSummary":null,"reviewSummaryHighlights":null,"reviewSummaryConcerns":null,"reviewSummaryGeneratedAt":null,"activationCount":0,"lastUsedAt":null,"lastSuccessfullyRanAt":null,"lastHealthCheckAt":"2026-09-15T12:34:04.989Z","isFirstParty":false}