{"slug":"mohitagw15856-llm-cost-latency-budget","source_name":"mohitagw15856/llm-cost-latency-budget","name":"Mohitagw15856/LLM Cost Latency Budget","description":"Model the cost and latency of an LLM feature before it ships and surprises the bill. Use when asked to estimate LLM API costs, set a latency/token budget, decide which model tier to use, or bring down the cost of an AI feature. Produces a cost & latency budget — token math per request, monthly cost projection, model tiering, caching/streaming levers, p95 latency targets, and a guardrail/alert plan.","version":1,"lift":{"pass_rate_delta_pts":27.27,"pass_rate_pct":100,"total_cases":22,"passed_cases":22,"tokens_delta_pct":47.3,"turns_delta_pct":0,"verdict":"pass","benchmark_model":"gemini-3.6-flash","grading_method":"judged","completed_at":"2026-08-15T14:44:03.841551+00:00"},"skill_score":1,"benchmark_models":[{"model":"gemini-3.6-flash","headline":true,"delta_pts":27.27,"with_pass_pct":100,"without_pass_pct":72.7,"tokens_delta_pct":47.3,"turns_delta_pct":0,"total_cases":22,"cases_aggregated":21,"verdict":"pass","never_hurt":true,"completed_at":"2026-08-15T14:44:03.841551+00:00","run_id":"76d2a7a8-0bd9-409c-988d-321d3a8df147","version_number":1,"is_latest_version":true,"gate":null}],"trust":{"skill_safety":"passed","safety_status":"clean","intent_verdict":"safe","content_status":"clean","indexable":true},"license":"MIT","install_count":0,"manifest_hash":"70552f28afb963f10403307b3f22fcc47bceffe26634eeb2cd8d530b0f8141cf","raw_url":"https://app.decimal.ai/s/mohitagw15856-llm-cost-latency-budget/SKILL.md","scorecard_url":"https://app.decimal.ai/skills/mohitagw15856-llm-cost-latency-budget"}