{"uid":"cap_fRtsFFQhqEVC9z-rK0poQ","slug":"gedx402-nemotron-3-120b-llm-via-x402-4c0f4fc0","name":"GEDX402 Nemotron-3-120B LLM via x402","description":"x402 workers ai. pay with usdc on base, polygon, arbitrum, world, or solana. no api keys.","url":"https://ged-x402-llm.jvalamis.workers.dev/v1/llm/nemotron-3-120b-a12b","method":"GET","headers":{},"bodySchema":{"type":"object","properties":{"messages":{"type":"array","items":{"type":"object"}},"max_tokens":{"type":"integer"}}},"responseSchema":{"type":"json","example":{"model":"@cf/nvidia/nemotron-3-120b-a12b","response":"Hello!"}},"example":null,"exampleRequest":null,"tags":["x402"],"displayCostAmount":"0.08","displayCostAsset":"USDC","priceDynamic":false,"priceHint":null,"priceStatus":"priced","priceSource":"probe","requiresHandshake":false,"reviewCount":0,"rating":{"score":"0.00","successRate":"0.00","reviews":0,"stars":null,"state":"unrated"},"availabilityStatus":"unknown","priceObserved":null,"sessionDeposit":null,"pricing":{"kind":"static","summary":"$0.08/call","primary":{"kind":"static","protocol":"x402","network":"base","amountUsd":"0.08","per":"call","confidence":"exact"},"accepted":[{"kind":"static","protocol":"x402","network":"base","amountUsd":"0.08","per":"call","confidence":"exact"}]},"paymentMethods":[{"uid":"pm_3sx6knhwgwMki6Tdrj2vl","protocol":"x402","methodType":"crypto","chain":"base","mode":"charge","costAmount":"0.08","costPer":"request","priority":0,"asset":"0x833589fCD6eDb6E08f4c7C32D4f71b54bdA02913","unit":"request","depositMicros":null,"planRef":null}],"brandName":null,"brandSlug":null,"brandBaseUrl":null,"brandDocsUrl":null,"whatItDoes":"Runs inference on NVIDIA Nemotron-3-120B-A12B via Cloudflare Workers AI, paid per-call with USDC over x402 protocol — no API keys required.","exampleAgentPrompt":"Ask the Nemotron-3-120B model on the x402 Workers AI endpoint: given the system message 'You are a helpful assistant' and the user message 'Explain quantum entanglement in simple terms', generate a reply with up to 512 tokens — I'll pay per call in USDC, no API key needed.","exampleUseCases":null,"resultDescription":"A JSON object containing the model identifier '@cf/nvidia/nemotron-3-120b-a12b' and a 'response' string with the assistant's generated text reply, up to the requested max_tokens.","failureModes":["Payment not received or insufficient USDC — x402 payment required before response is returned","Invalid message role — only 'system', 'user', 'assistant' are accepted","max_tokens exceeds 4096 — request rejected with validation error","Missing required 'messages' array — returns 400 bad request","Network or Workers AI backend error — upstream 5xx response","Unsupported payment network — only Base, Polygon, Arbitrum, World, Solana accepted"],"whenToPreferThis":"Choose this endpoint when you need a very large (120B parameter) NVIDIA Nemotron model for high-quality text generation and want to pay per call in USDC via x402 without managing API keys or subscriptions. Prefer it over smaller models when response quality matters more than cost, and over traditional LLM APIs when you want frictionless crypto-native access.","instructions":null,"reviewSummary":null,"reviewSummaryHighlights":null,"reviewSummaryConcerns":null,"reviewSummaryGeneratedAt":null,"activationCount":0,"lastUsedAt":null,"lastSuccessfullyRanAt":null,"lastHealthCheckAt":"2026-09-14T18:30:30.563Z","isFirstParty":false}