{
  "slug": "read-results",
  "title": "Read a coding-agent result",
  "lead": "Start with the task, then each measure.",
  "benefit": "Code-season results are pending. No leader is established on any measure.",
  "visual": [
    "Task scope",
    "Measure",
    "Evidence"
  ],
  "visualLabel": "Reading order",
  "cta": {
    "label": "Review the reading method",
    "href": "/read-results#measures"
  },
  "sections": [
    {
      "id": "scope",
      "title": "Check whether the task fits",
      "text": "Small Python tasks; production relevance is unproven.",
      "details": [
        {
          "title": "Read the scope and limits",
          "text": "Season one concerns small, original Python repositories. Forty pilot tasks precede a separate official corpus. It does not establish performance on production codebases."
        }
      ]
    },
    {
      "id": "measures",
      "title": "Read each measure separately",
      "text": "A measure is not an overall ranking.",
      "details": [
        {
          "title": "Read the scope and limits",
          "text": "Read quality, coverage, cost and speed with their evidence and uncertainty. A lead on one measure does not establish an overall champion. An unmeasured dimension stays unmeasured."
        }
      ]
    },
    {
      "id": "consideration",
      "title": "Look for missing alternatives",
      "text": "Observed alternatives, not admitted participants.",
      "details": [
        {
          "title": "Read the scope and limits",
          "text": "Two independent organizations mention five options: Claude Code, Codex CLI, GitHub Copilot coding agent, Cursor and Gemini CLI. Claude Code, Codex CLI and Cursor appear in both accounts. Mentions are not performance results."
        },
        {
          "title": "No coverage claim follows yet",
          "text": "This observed consideration set was not frozen before the rights review. It is a candidate, not an admitted roster or the denominator of a published coverage claim. The Copilot account concerns its cloud coding agent, not local CLI adoption."
        }
      ]
    },
    {
      "id": "evidence",
      "title": "Follow the evidence",
      "text": "Keep each claim within its evidence.",
      "details": [
        {
          "title": "Read the scope and limits",
          "text": "Check configurations, dates, corpus, frozen rules and missing evidence. Historical examples belong to their own categories. They do not establish a coding-agent result."
        },
        {
          "title": "Reuse a published result",
          "text": "For a published result, keep its supplied citation, license and attribution together. The citation identifies the run, observation date, status, scope and canonical result URL. Preserve uncertainty and missing evidence when reusing a claim. If this metadata is absent, do not invent it or treat the record as ready for publication. CC-BY-4.0 covers published TORNEO measurements and result tables; third-party archives retain their original rights. Citation metadata does not establish publication, scientific validity or independent review."
        }
      ]
    }
  ],
  "schema": "torneo.marketing-page.v1",
  "version": "code-season.3",
  "consideration": {
    "status": "OBSERVED_CANDIDATE",
    "freeze_sha256": "9fdc4f5da8f9e31baffb12e04faebb24c869f9aafcecd75b34d572bab9f9f0ff",
    "E": [
      "claudecode",
      "codexcli",
      "copilotagent",
      "cursor",
      "geminicli"
    ],
    "K": [
      "claudecode",
      "codexcli",
      "cursor"
    ],
    "names": {
      "claudecode": "Claude Code",
      "codexcli": "Codex CLI",
      "copilotagent": "GitHub Copilot coding agent",
      "cursor": "Cursor",
      "geminicli": "Gemini CLI"
    }
  },
  "url": "https://torneo.ai/read-results"
}
