{"uid":"cap_vN7vJmvSJTZT05nfAQjrw","slug":"x402-nvidia-llama-3-3-nemotron-super-49b-chat-completions-8998a2ed","name":"X402-NVIDIA Llama-3.3-Nemotron-Super-49B Chat Completions","description":"Access 48+ NVIDIA NIM AI models via x402 micropayments. Chat completions, vision, safety, translation, and more.","url":"https://x402-nvidia.vercel.app/api/llama-3-3-nemotron-super-49b","method":"POST","headers":{},"bodySchema":null,"responseSchema":{"type":"object","example":{"choices":[{"message":{"role":"assistant","content":"Hello!"}}]},"properties":{"choices":{"type":"array","description":"Model responses"}}},"example":null,"exampleRequest":null,"tags":["x402"],"displayCostAmount":"0.158628","displayCostAsset":"USDC","priceDynamic":false,"priceHint":null,"priceStatus":"priced","priceSource":"probe","requiresHandshake":false,"reviewCount":0,"rating":{"score":"0.00","successRate":"0.00","reviews":0,"stars":null,"state":"unrated"},"availabilityStatus":"unknown","priceObserved":null,"sessionDeposit":null,"pricing":{"kind":"static","summary":"$0.158628/call","primary":{"kind":"static","protocol":"x402","network":"base","amountUsd":"0.158628","per":"call","confidence":"exact"},"accepted":[{"kind":"static","protocol":"x402","network":"base","amountUsd":"0.158628","per":"call","confidence":"exact"}]},"paymentMethods":[{"uid":"pm_WIBhsGuBuw8qV9E7-13bd","protocol":"x402","methodType":"crypto","chain":"base","mode":"charge","costAmount":"0.158628","costPer":"request","priority":0,"asset":"0x833589fCD6eDb6E08f4c7C32D4f71b54bdA02913","unit":"request","depositMicros":null,"planRef":null}],"brandName":null,"brandSlug":null,"brandBaseUrl":null,"brandDocsUrl":null,"whatItDoes":"Run chat completions against NVIDIA's Llama-3.3-Nemotron-Super-49B model via x402 micropayment-gated API","exampleAgentPrompt":"Use NVIDIA's Llama-3.3-Nemotron-Super-49B model to answer this question: 'What are the key differences between transformer and mamba architectures for long-context tasks?'","exampleUseCases":null,"resultDescription":"Returns a chat completion object with an array of choices, each containing a message with role 'assistant' and the generated text content from the Nemotron-Super-49B model.","failureModes":["Payment failure: x402 micropayment rejected or insufficient USDC balance","Model overload: NVIDIA NIM backend unavailable or rate-limited","Malformed input: missing or invalid message format causes 400 error","Network timeout: Vercel proxy or NVIDIA NIM endpoint unreachable","Authentication error: x402 payment header missing or invalid"],"whenToPreferThis":"Choose this endpoint when you need access to NVIDIA's Llama-3.3-Nemotron-Super-49B model specifically, want pay-per-call pricing via USDC micropayments without a subscription, or need a high-parameter (49B) model for complex reasoning tasks billed per request.","instructions":null,"reviewSummary":null,"reviewSummaryHighlights":null,"reviewSummaryConcerns":null,"reviewSummaryGeneratedAt":null,"activationCount":0,"lastUsedAt":null,"lastSuccessfullyRanAt":null,"lastHealthCheckAt":"2026-09-14T06:41:38.945Z","isFirstParty":false}