{"uid":"cap_IffbiUAw7nstF_1mSz71N","slug":"nvidia-nemotron-3-ultra-550b-via-x402-micropayments-c2a80acb","name":"NVIDIA Nemotron-3 Ultra 550B via x402 Micropayments","description":"Access 48+ NVIDIA NIM AI models via x402 micropayments. Chat completions, vision, safety, translation, and more.","url":"https://x402-nvidia.vercel.app/api/nemotron-3-ultra-550b","method":"POST","headers":{},"bodySchema":null,"responseSchema":{"type":"object","example":{"choices":[{"message":{"role":"assistant","content":"Hello!"}}]},"properties":{"choices":{"type":"array","description":"Model responses"}}},"example":null,"exampleRequest":null,"tags":["x402"],"displayCostAmount":"0.129626","displayCostAsset":"USDC","priceDynamic":false,"priceHint":null,"priceStatus":"priced","priceSource":"probe","requiresHandshake":false,"reviewCount":0,"rating":{"score":"0.00","successRate":"0.00","reviews":0,"stars":null,"state":"unrated"},"availabilityStatus":"unknown","priceObserved":null,"sessionDeposit":null,"pricing":{"kind":"static","summary":"$0.129626/call","primary":{"kind":"static","protocol":"x402","network":"base","amountUsd":"0.129626","per":"call","confidence":"exact"},"accepted":[{"kind":"static","protocol":"x402","network":"base","amountUsd":"0.129626","per":"call","confidence":"exact"}]},"paymentMethods":[{"uid":"pm_Fz_g7EH6dz67fzntDN2-K","protocol":"x402","methodType":"crypto","chain":"base","mode":"charge","costAmount":"0.129626","costPer":"request","priority":0,"asset":"0x833589fCD6eDb6E08f4c7C32D4f71b54bdA02913","unit":"request","depositMicros":null,"planRef":null}],"brandName":null,"brandSlug":null,"brandBaseUrl":null,"brandDocsUrl":null,"whatItDoes":"Run inference on NVIDIA's Nemotron-3 Ultra 550B large language model via pay-per-call USDC micropayments using the x402 protocol","exampleAgentPrompt":"Use the NVIDIA Nemotron-3 Ultra 550B model to answer this question with detailed reasoning: 'What are the key tradeoffs between transformer and state-space architectures for long-context language modeling?'","exampleUseCases":null,"resultDescription":"Returns a chat completion object with an array of choices, each containing an assistant message with the model's generated response content. Structured as a standard OpenAI-compatible chat completion response.","failureModes":["Payment failure if x402 USDC micropayment is rejected or wallet has insufficient funds","402 Payment Required if x402 header is missing or malformed","Model unavailable or timeout for very long prompts against the 550B model","Malformed request body returns 400 if messages array is missing or improperly structured","Rate limiting if too many concurrent requests are sent"],"whenToPreferThis":"Choose this endpoint when you need a very large (550B parameter) NVIDIA language model for complex reasoning, nuanced generation, or high-quality completions, and want to pay per call in USDC via x402 without a subscription. Prefer over smaller models when task complexity demands frontier-scale reasoning capacity.","instructions":null,"reviewSummary":null,"reviewSummaryHighlights":null,"reviewSummaryConcerns":null,"reviewSummaryGeneratedAt":null,"activationCount":0,"lastUsedAt":null,"lastSuccessfullyRanAt":null,"lastHealthCheckAt":"2026-09-13T12:33:49.240Z","isFirstParty":false}