{"uid":"cap_bND430fl2KWngVKhW_YZk","slug":"modelprices-xyz-google-vertex-ai-pricing-table-f8f3f6e9","name":"modelprices.xyz Google Vertex AI Pricing Table","description":"Google Vertex AI pricing table: per-token cost of every AI model Google Vertex AI serves, in one call — input, output, cache and batch USD per 1M tokens, ranked cheapest first, with context window and capability flags joined in. Compare LLM token cost inside Google Vertex AI and against other providers hosting the same model. Normalized from public sources, cross-checked, refreshed hourly.","url":"https://modelprices.xyz/llm/prices/vertex-ai","method":"GET","headers":{},"bodySchema":{"type":"object","$schema":"https://json-schema.org/draft/2020-12/schema","required":["input"],"properties":{"input":{"type":"object","required":["type","method"],"properties":{"type":{"type":"string","const":"http"},"method":{"enum":["GET"],"type":"string"},"queryParams":{"type":"object","properties":{}}},"additionalProperties":false},"output":{"type":"object","required":["type"],"properties":{"type":{"type":"string"},"example":{"type":"object"}}}}},"responseSchema":null,"example":null,"exampleRequest":null,"tags":["x402"],"displayCostAmount":"0.01","displayCostAsset":"USDC","priceDynamic":false,"priceHint":null,"priceStatus":"priced","priceSource":"probe","requiresHandshake":false,"reviewCount":0,"rating":{"score":"0.00","successRate":"0.00","reviews":0,"stars":null,"state":"unrated"},"availabilityStatus":"unknown","priceObserved":null,"sessionDeposit":null,"pricing":{"kind":"static","summary":"$0.01/call","primary":{"kind":"static","protocol":"x402","network":"base","amountUsd":"0.01","per":"call","confidence":"exact"},"accepted":[{"kind":"static","protocol":"x402","network":"base","amountUsd":"0.01","per":"call","confidence":"exact"}]},"paymentMethods":[{"uid":"pm_UILHMhO9QS8XHXeibtFRI","protocol":"x402","methodType":"crypto","chain":"base","mode":"charge","costAmount":"0.01","costPer":"request","priority":0,"asset":"0x833589fCD6eDb6E08f4c7C32D4f71b54bdA02913","unit":"request","depositMicros":null,"planRef":null}],"brandName":null,"brandSlug":null,"brandBaseUrl":null,"brandDocsUrl":null,"whatItDoes":"Returns per-token pricing for every AI model available on Google Vertex AI, including input, output, cache, and batch costs in USD per 1M tokens, ranked cheapest first with context window and capability metadata.","exampleAgentPrompt":"Pull the full Google Vertex AI pricing table so I can see what every model costs per million input and output tokens, ranked cheapest first, and check if any support caching or batch discounts.","exampleUseCases":[{"title":"Cost-optimize LLM provider selection","prompt":"I'm building an app that calls Vertex AI a lot — can you get me the full Vertex AI pricing table ranked cheapest first so I can pick the most affordable model that fits my context window needs?"},{"title":"Cross-provider model cost comparison","prompt":"I want to compare what Gemini 1.5 Pro costs on Vertex AI versus other hosting providers — can you fetch the Vertex AI pricing data so I can see its per-token rates for input, output, and batch?"},{"title":"Budget forecasting for AI workloads","prompt":"We're forecasting our monthly AI spend — can you get the current Vertex AI token prices for all their models including any cache and batch rates so I can plug them into our cost model?"}],"resultDescription":"A ranked list of all AI models available on Google Vertex AI, each with USD per 1M token prices for input, output, cache read, and batch inference, plus context window size in tokens and capability flags (e.g. vision, reasoning), sorted ascending by input cost. Data is normalized from public sources and refreshed hourly.","failureModes":["Upstream Vertex AI pricing page unavailable — stale cached data returned or 503","Model not yet indexed — newly released models may have a lag before appearing","Rate limiting if called excessively — 402 or 429 response","Pricing data mismatch if Google updates rates between hourly refreshes"],"whenToPreferThis":"Use this endpoint when you need a comprehensive, pre-normalized, machine-readable Vertex AI pricing table in a single call — especially for cost comparisons across models or providers. Prefer this over scraping Google's pricing pages directly, as it delivers ranked, cross-checked data with context window and capability metadata already joined. Best for agents automating LLM cost analysis, budget forecasting, or dynamic model selection within the Vertex AI ecosystem.","instructions":null,"reviewSummary":null,"reviewSummaryHighlights":null,"reviewSummaryConcerns":null,"reviewSummaryGeneratedAt":null,"activationCount":0,"lastUsedAt":null,"lastSuccessfullyRanAt":null,"lastHealthCheckAt":"2026-09-14T18:44:14.921Z","isFirstParty":false}