{"uid":"cap_8ErCE6jz2OYQ5Z5CVhE0r","slug":"arbipulse-ai-inference-price-arbitrage-7c7c17ef","name":"ArbiPulse AI Inference Price Arbitrage","description":"AI inference price arbitrage — every provider serving a given model with live input/output $/Mtoken, blended 3:1 cost, cheapest-vs-dearest spread multiple, quantization caveats. Same open-weights model spans 2-5x across providers; for agents that buy inference this is self-referential spend recovery. Live listings, no LLM.","url":"https://arbipulse.theaslangroupllc.com/api/inference-arb","method":"GET","headers":{},"bodySchema":{"type":"object","$schema":"https://json-schema.org/draft/2020-12/schema","required":["input"],"properties":{"input":{"type":"object","required":["type","method"],"properties":{"type":{"type":"string","const":"http"},"method":{"enum":["GET","HEAD","DELETE"],"type":"string"},"queryParams":{"type":"object","properties":{"model":{"type":"string"}}}},"additionalProperties":false},"output":{"type":"object","required":["type"],"properties":{"type":{"type":"string"},"errors":{"type":"object","description":"Documented error responses, keyed by HTTP status code","additionalProperties":{"type":"object","required":["description"],"properties":{"example":{"type":"object"},"description":{"type":"string"}}}},"example":{"type":"object"}}}}},"responseSchema":{"type":"json","example":{"model":"meta-llama/llama-3.3-70b-instruct","cheapest":"DeepInfra","providers":[{"provider":"DeepInfra","quantization":"fp8","blended_usd_per_mtok":0.2}],"spread_multiple":4.6}},"example":null,"exampleRequest":null,"tags":["x402"],"displayCostAmount":"0.05","displayCostAsset":"USDC","priceDynamic":false,"priceHint":null,"priceStatus":"priced","priceSource":"probe","requiresHandshake":false,"reviewCount":0,"rating":{"score":"0.00","successRate":"0.00","reviews":0,"stars":null,"state":"unrated"},"availabilityStatus":"unknown","priceObserved":null,"sessionDeposit":null,"pricing":{"kind":"static","summary":"$0.05/call","primary":{"kind":"static","protocol":"x402","network":"base","amountUsd":"0.05","per":"call","confidence":"exact"},"accepted":[{"kind":"static","protocol":"x402","network":"base","amountUsd":"0.05","per":"call","confidence":"exact"}]},"paymentMethods":[{"uid":"pm_aaj4TPlBkk9lHN707Ogfi","protocol":"x402","methodType":"crypto","chain":"base","mode":"charge","costAmount":"0.05","costPer":"request","priority":0,"asset":"0x833589fCD6eDb6E08f4c7C32D4f71b54bdA02913","unit":"request","depositMicros":null,"planRef":null}],"brandName":null,"brandSlug":null,"brandBaseUrl":null,"brandDocsUrl":null,"whatItDoes":"Returns live pricing for every provider serving a given open-weights LLM model, including input/output cost per million tokens, blended 3:1 cost, cheapest-to-dearest spread multiple, and quantization caveats.","exampleAgentPrompt":"Show me all providers currently serving Llama 3 70B, with their input and output prices per million tokens, the blended 3:1 cost, the cheapest-to-dearest spread multiple, and any quantization caveats — I want to pick the cheapest option for my agent's inference spend.","exampleUseCases":null,"resultDescription":"A live, structured listing of every provider offering the requested open-weights model, each entry including: input $/Mtoken, output $/Mtoken, blended 3:1 cost, the spread multiple between the cheapest and most expensive provider, and any quantization or precision caveats. No LLM-generated content — data is pulled live.","failureModes":["Model name not recognized or not tracked — returns empty provider list or 404-style error","Model has only one known provider — spread multiple is 1x, caveat returned","Pricing data temporarily stale — endpoint may return a freshness warning","Payment not processed — x402 payment required error before data is returned"],"whenToPreferThis":"Use this endpoint when an AI agent needs to minimize its own inference spend on a specific open-weights model, when comparing provider pricing before routing workloads, or when auditing inference cost arbitrage opportunities. Prefer over manual provider docs because it aggregates live pricing across 2-5+ providers simultaneously with a pre-computed spread multiple and blended cost metric.","instructions":null,"reviewSummary":null,"reviewSummaryHighlights":null,"reviewSummaryConcerns":null,"reviewSummaryGeneratedAt":null,"activationCount":0,"lastUsedAt":null,"lastSuccessfullyRanAt":null,"lastHealthCheckAt":"2026-09-14T18:33:48.451Z","isFirstParty":false}