{"slug":"agentsope-agentsop-code-execution-decision","source_name":"agentsope/agentsop-code-execution-decision","name":"Agentsope/Agentsop Code Execution Decision","description":"Decision rubric for when an LM agent should write-and-run code (Program-of-Thought / code interpreter) versus reason in natural language: classify each step as deterministic- computable (emit + execute code, feed the result back) vs judgment (stay in prose). Use when designing or debugging an agent step that does arithmetic/parsing/data transforms, when prose reasoning hallucinates a computation (under-coding), or when a sandbox round- trip is wasted on a judgment task (over-coding). Search keyw","version":1,"lift":{"pass_rate_delta_pts":13.04,"pass_rate_pct":95.7,"total_cases":23,"passed_cases":22,"tokens_delta_pct":300.4,"turns_delta_pct":0,"verdict":"mixed","benchmark_model":"gemini-3.6-flash","grading_method":"judged","completed_at":"2026-08-24T13:47:33.298513+00:00"},"skill_score":0.9565,"benchmark_models":[{"model":"gemini-3.6-flash","headline":true,"delta_pts":13.04,"with_pass_pct":95.7,"without_pass_pct":82.6,"tokens_delta_pct":300.4,"turns_delta_pct":0,"total_cases":23,"cases_aggregated":23,"verdict":"mixed","never_hurt":true,"completed_at":"2026-08-24T13:47:33.298513+00:00","run_id":"87a9e64c-7a9c-44f5-92b0-4a68ffcd71a7","version_number":1,"is_latest_version":true,"gate":null}],"trust":{"skill_safety":"passed","safety_status":"clean","intent_verdict":"safe","content_status":"clean","indexable":true},"license":"MIT","install_count":0,"manifest_hash":"ca2961e7da060ef8bbff116582903330ad43d502cf883d9006b507767625811e","raw_url":"https://app.decimal.ai/s/agentsope-agentsop-code-execution-decision/SKILL.md","scorecard_url":"https://app.decimal.ai/skills/agentsope-agentsop-code-execution-decision"}