{"uid":"cap_Q42Afl7bkoyaEHBzVdB1P","slug":"x402engine-deepseek-v4-flash-llm-8efee2b0","name":"x402engine DeepSeek V4 Flash LLM","description":"DeepSeek's low-latency V4 model — ultra-low-cost reasoning and coding with 1M context","url":"https://x402engine.app/api/llm/deepseek-v4-flash?utm_source=zero.xyz","method":"POST","headers":{},"bodySchema":null,"responseSchema":null,"example":null,"exampleRequest":null,"tags":["x402"],"displayCostAmount":"0.003","displayCostAsset":"USDC","priceDynamic":false,"priceHint":null,"priceStatus":"priced","priceSource":"registry","requiresHandshake":false,"reviewCount":0,"rating":{"score":"0.00","successRate":"0.00","reviews":0,"stars":null,"state":"unrated"},"availabilityStatus":"unknown","priceObserved":null,"sessionDeposit":null,"pricing":{"kind":"static","summary":"$0.003/call","primary":{"kind":"static","protocol":"x402","network":"base","amountUsd":"0.003","per":"call","confidence":"exact"},"accepted":[{"kind":"static","protocol":"x402","network":"base","amountUsd":"0.003","per":"call","confidence":"exact"}]},"paymentMethods":[{"uid":"pm_LlaBcC0JFpIdhJ0x8J-yp","protocol":"x402","methodType":"crypto","chain":"base","mode":"charge","costAmount":"0.003","costPer":"request","priority":0,"asset":null,"unit":"request","depositMicros":null,"planRef":null}],"brandName":null,"brandSlug":null,"brandBaseUrl":null,"brandDocsUrl":null,"whatItDoes":"Runs DeepSeek's low-latency V4 Flash model for ultra-low-cost reasoning and coding tasks with up to 1M token context, accessible via x402 micropayment at $0.003 USDC per call","exampleAgentPrompt":"Using the DeepSeek V4 Flash model, write me a Python function that parses a JSON log file and extracts all error messages with their timestamps — keep it concise and well-commented.","exampleUseCases":[{"title":"Automated code review for PRs","prompt":"Review this Python pull request diff and point out any bugs, edge cases, or style issues — use DeepSeek V4 Flash so it's fast and cheap: [paste diff here]"},{"title":"Long document summarization","prompt":"Summarize this 200-page legal contract into a concise executive summary hitting the key obligations, deadlines, and risk clauses — use a 1M context model like DeepSeek V4 Flash to handle the full text."},{"title":"Agentic reasoning task on a budget","prompt":"I need to break down a complex multi-step data pipeline problem into subtasks and propose a solution architecture — use DeepSeek V4 Flash to reason through it without burning a lot of API budget."}],"resultDescription":"Returns a text completion from the DeepSeek V4 Flash model, including the generated content (code, reasoning, prose, etc.) and token usage metadata. Response follows a standard LLM chat completion format.","failureModes":["402 Payment Required if x402 USDC micropayment header is missing or insufficient","prompt exceeds 1M token context window resulting in truncation or error","rate limiting if too many concurrent requests are sent","model returns truncated output if max_tokens is set too low","malformed request body causes 400 error"],"whenToPreferThis":"Choose this endpoint when you need a fast, very low-cost LLM inference call (especially for coding or reasoning) and want to pay per-call via USDC micropayments rather than a subscription. Ideal for agentic workflows that invoke LLMs frequently at scale, or when you need a 1M token context window without committing to a monthly plan. Prefer over GPT-4 or Claude when cost-per-call is the primary constraint and DeepSeek-quality output is sufficient.","instructions":null,"reviewSummary":null,"reviewSummaryHighlights":null,"reviewSummaryConcerns":null,"reviewSummaryGeneratedAt":null,"activationCount":0,"lastUsedAt":null,"lastSuccessfullyRanAt":null,"lastHealthCheckAt":"2026-10-02T00:32:54.398Z","isFirstParty":false,"canonicalSlug":"x402engine-deepseek-v4-flash-llm-8efee2b0"}