{"uid":"cap_PB3Qa6oWGTGT3lB0rkdad","slug":"pqs-prompt-quality-score-cross-model-comparison-02ad9624","name":"PQS: Prompt Quality Score – Cross-Model Comparison","description":"Cross-model scoring: Claude Sonnet 4 vs GPT-4o, judged by a third model","url":"https://api.relai.fi/relay/1779373065272/api/score/compare","method":"GET","headers":{},"bodySchema":{"type":"object","properties":{"prompt":{"type":"string","maxLength":10000,"minLength":1,"description":"Prompt to run through both Claude and GPT-4o for cross-model comparison"},"vertical":{"enum":["software","content","business","education","science","crypto","general","research"],"type":"string","default":"general","description":"Domain: software/content/business/education/science/crypto/general/research"}}},"responseSchema":null,"example":null,"exampleRequest":null,"tags":["x402"],"displayCostAmount":"1.25","displayCostAsset":"USDC","priceDynamic":false,"priceHint":null,"priceStatus":"priced","priceSource":"probe","requiresHandshake":false,"reviewCount":0,"rating":{"score":"0.00","successRate":"0.00","reviews":0,"stars":null,"state":"unrated"},"availabilityStatus":"unknown","priceObserved":null,"sessionDeposit":null,"pricing":{"kind":"static","summary":"$1.25/call","primary":{"kind":"static","protocol":"x402","network":"base","amountUsd":"1.25","per":"call","confidence":"exact"},"accepted":[{"kind":"static","protocol":"x402","network":"base","amountUsd":"1.25","per":"call","confidence":"exact"}]},"paymentMethods":[{"uid":"pm_AiMdBBowceoRNFPnXgVf5","protocol":"x402","methodType":"crypto","chain":"base","mode":"charge","costAmount":"1.25","costPer":"request","priority":0,"asset":null,"unit":"request","depositMicros":null,"planRef":null}],"brandName":null,"brandSlug":null,"brandBaseUrl":null,"brandDocsUrl":null,"whatItDoes":"Scores a prompt by running it against Claude Sonnet 4 and GPT-4o in parallel, then uses a third model to judge and compare the output quality","exampleAgentPrompt":"Can you score the quality of this prompt using PQS — run it against both Claude Sonnet 4 and GPT-4o and tell me which model handles it better, using the 'software' vertical: 'You are a senior engineer. Review the following Python function for correctness and suggest improvements.'","exampleUseCases":null,"resultDescription":"A structured quality score comparing how Claude Sonnet 4 and GPT-4o responded to the given prompt, judged by a third model, with per-model ratings and an overall verdict on which model performed better in the specified domain vertical.","failureModes":["Prompt text missing or empty — returns 400 validation error","Prompt exceeds 10,000 character limit — returns 400","Invalid vertical enum value — returns 400 schema validation error","Payment not included or insufficient — returns 402 Payment Required","Upstream model API timeout — may return 503 or partial result","Judge model failure — scoring may be incomplete or unavailable"],"whenToPreferThis":"Choose this endpoint when you need an objective, third-party judgment of prompt quality across multiple leading LLMs simultaneously. It is ideal for prompt engineers, AI developers, and researchers who want to benchmark their prompts against both Claude Sonnet 4 and GPT-4o rather than testing each model manually. Prefer it over single-model evaluation when cross-model comparison or domain-specific scoring (software, research, crypto, etc.) is required.","instructions":null,"reviewSummary":null,"reviewSummaryHighlights":null,"reviewSummaryConcerns":null,"reviewSummaryGeneratedAt":null,"activationCount":0,"lastUsedAt":null,"lastSuccessfullyRanAt":null,"lastHealthCheckAt":"2026-09-14T06:43:47.067Z","isFirstParty":false}