{
  "benchmark": "Warehouse shift v1 \u2014 cost addendum",
  "date": "2026-09-30",
  "suite_sha256": "0770cc0771a4a6d8db38823ba6282d65f5a1bf589f045a1f35fc42a1203e7094",
  "metric": "Mean estimated model USD for the three primary scored shifts; includes failed checks; overhead reported separately",
  "excluded_costs": [
    "OmniLink subscription",
    "Codex subscription or credits",
    "Claude subscription",
    "hosting",
    "local compute",
    "tax"
  ],
  "codex_estimate_conditions": {
    "model": "gpt-6.1-sol",
    "reasoning_effort": "high",
    "service_tier_assumption": "Standard API, no Fast/Ultrafast or regional premium",
    "rates_usd_per_million": {
      "input": 2.0,
      "cached_input": 0.1,
      "cache_write": 2.5,
      "output": 10.0
    },
    "pricing_checked": "2026-09-30",
    "long_context_threshold": 272000,
    "observed_cache_write_tokens": 0,
    "accounting": "Latest native cumulative totals per thread, not a sum of repeated usage notifications. Cached and cache-write counts are subsets of input; reasoning is included once in output. One primary interrupted turn had no individual usage event; its billed usage is unknown. Native counters are best-effort, not an invoice. Zero recorded native cache writes do not establish API write charges. Historical billing tier was not preserved; current config cannot establish past billing. These API rates do not estimate included subscription usage. Development and this coding chat's work are excluded.",
    "sources": [
      "https://developers.openai.com/api/docs/models/gpt-6.1-sol",
      "https://developers.openai.com/api/docs/guides/prompt-caching",
      "https://developers.openai.com/api/docs/guides/agents-api/observability",
      "https://learn.chatgpt.com/docs/pricing"
    ]
  },
  "original_rates_usd_per_million": {
    "input": 1.5,
    "cached_input": 0.15,
    "output_including_thinking": 9.0
  },
  "source_archive_sha256": {
    "evidence.zip": "724e05061bad7a5e06ba353080673b13262e1ad6327a6a13a27748d1ba329cca",
    "codex-evidence.zip": "9cfbedcd3d0981c1d0a55003243bb9f7e28987ebd9e5fad59a65713d7cf62186",
    "claude-evidence.zip": "c66413a79c470f057a9ca439ee75986769dc32fc60aa067ac8e17f79d6537e7b"
  },
  "arms": {
    "omnilink": {
      "model": "gemini-3.5-flash",
      "scored_shifts": 3,
      "scored_usd": 2.709942,
      "usd_per_scored_shift": 0.903314,
      "overhead_usd": 0.0,
      "all_attempts_usd": 2.709942,
      "comparable_cost": true,
      "attempts": 3,
      "dollar_basis": "Recorded usage at original provider list rates"
    },
    "plain_basic": {
      "model": "gemini-3.5-flash",
      "scored_shifts": 3,
      "scored_usd": 4.255651,
      "usd_per_scored_shift": 1.4185503333333334,
      "overhead_usd": 0.7558150000000001,
      "all_attempts_usd": 5.011466,
      "comparable_cost": true,
      "attempts": 4,
      "dollar_basis": "Recorded usage at original provider list rates"
    },
    "langgraph_basic": {
      "model": "gemini-3.5-flash",
      "scored_shifts": 3,
      "scored_usd": 4.032869,
      "usd_per_scored_shift": 1.3442896666666666,
      "overhead_usd": 1.9033439999999997,
      "all_attempts_usd": 5.9362129999999995,
      "comparable_cost": true,
      "attempts": 4,
      "dollar_basis": "Recorded usage at original provider list rates"
    },
    "lobster_basic": {
      "model": "gemini-3.5-flash",
      "scored_shifts": 3,
      "scored_usd": 4.001338,
      "usd_per_scored_shift": 1.3337793333333332,
      "overhead_usd": 0.9893550000000007,
      "all_attempts_usd": 4.990693,
      "comparable_cost": true,
      "attempts": 4,
      "dollar_basis": "Recorded usage at original provider list rates"
    },
    "plain_full": {
      "model": "gemini-3.5-flash",
      "scored_shifts": 3,
      "scored_usd": 6.549084,
      "usd_per_scored_shift": 2.1830279999999997,
      "overhead_usd": 1.585667,
      "all_attempts_usd": 8.134751,
      "comparable_cost": true,
      "attempts": 4,
      "dollar_basis": "Recorded usage at original provider list rates"
    },
    "langgraph_full": {
      "model": "gemini-3.5-flash",
      "scored_shifts": 3,
      "scored_usd": 6.583503,
      "usd_per_scored_shift": 2.1945010000000003,
      "overhead_usd": 2.1632099999999994,
      "all_attempts_usd": 8.746713,
      "comparable_cost": true,
      "attempts": 4,
      "dollar_basis": "Recorded usage at original provider list rates"
    },
    "lobster_full": {
      "model": "gemini-3.5-flash",
      "scored_shifts": 3,
      "scored_usd": 5.27051,
      "usd_per_scored_shift": 1.7568366666666666,
      "overhead_usd": 3.803979,
      "all_attempts_usd": 9.074489,
      "comparable_cost": true,
      "attempts": 5,
      "dollar_basis": "Recorded usage at original provider list rates"
    },
    "codex_full": {
      "model": "gpt-6.1-sol",
      "scored_shifts": 3,
      "scored_usd": 2.7378160000000005,
      "usd_per_scored_shift": 0.9126053333333335,
      "overhead_usd": 2.3725788,
      "all_attempts_usd": 5.1103948,
      "attempts": 6,
      "comparable_cost": false,
      "dollar_basis": "Conditional recorded-token Standard API pricing scenario; excluded from cost ranking; actual signed-in cost unavailable",
      "billed_usd": null
    },
    "claude_full": {
      "model": "claude-opus-5-5",
      "scored_shifts": 3,
      "attempts": 3,
      "scored_usd": null,
      "usd_per_scored_shift": null,
      "overhead_usd": null,
      "all_attempts_usd": null,
      "comparable_cost": false,
      "billed_usd": null,
      "cli_list_price_estimate_usd": 3.7243320000000004,
      "cli_list_price_estimate_per_shift_usd": 1.2414440000000002,
      "dollar_basis": "Claude Code's own API list-price estimate, not a bill; actual signed-in cost unavailable; excluded from ranking"
    }
  },
  "codex_attempts": [
    {
      "run": "codex-shift-v1-20260930-corrected-r0",
      "primary": true,
      "outcome": "FAIL",
      "tokens": {
        "totalTokens": 5212373,
        "inputTokens": 5205883,
        "cachedInputTokens": 5090688,
        "cacheWriteInputTokens": 0,
        "outputTokens": 6490,
        "reasoningOutputTokens": 1840
      },
      "max_observed_request_input_tokens": 75508,
      "elapsed_s": 1872.9679999999935,
      "turns_without_individual_usage": 0,
      "interrupted_turns_without_individual_usage": 0,
      "standard_api_equivalent_usd": 0.8043588
    },
    {
      "run": "codex-shift-v1-20260930-corrected-r1",
      "primary": true,
      "outcome": "FAIL",
      "tokens": {
        "totalTokens": 6359936,
        "inputTokens": 6352886,
        "cachedInputTokens": 6248320,
        "cacheWriteInputTokens": 0,
        "outputTokens": 7050,
        "reasoningOutputTokens": 1855
      },
      "max_observed_request_input_tokens": 84129,
      "elapsed_s": 1869.4380000000237,
      "turns_without_individual_usage": 1,
      "interrupted_turns_without_individual_usage": 1,
      "standard_api_equivalent_usd": 0.904464
    },
    {
      "run": "codex-shift-v1-20260930-corrected-r2",
      "primary": true,
      "outcome": "FAIL",
      "tokens": {
        "totalTokens": 7233063,
        "inputTokens": 7226903,
        "cachedInputTokens": 7098112,
        "cacheWriteInputTokens": 0,
        "outputTokens": 6160,
        "reasoningOutputTokens": 1479
      },
      "max_observed_request_input_tokens": 108280,
      "elapsed_s": 1876.3600000000442,
      "turns_without_individual_usage": 0,
      "interrupted_turns_without_individual_usage": 0,
      "standard_api_equivalent_usd": 1.0289932000000002
    },
    {
      "run": "codex-shift-v1-20260930-r0",
      "primary": false,
      "outcome": "FAIL",
      "tokens": {
        "totalTokens": 5080144,
        "inputTokens": 5074259,
        "cachedInputTokens": 4919168,
        "cacheWriteInputTokens": 0,
        "outputTokens": 5885,
        "reasoningOutputTokens": 1557
      },
      "max_observed_request_input_tokens": 75471,
      "elapsed_s": 1871.0779999999795,
      "turns_without_individual_usage": 1,
      "interrupted_turns_without_individual_usage": 1,
      "standard_api_equivalent_usd": 0.8609488000000001
    },
    {
      "run": "codex-shift-v1-20260930-r1",
      "primary": false,
      "outcome": "ERROR",
      "tokens": {
        "totalTokens": 5068212,
        "inputTokens": 5061748,
        "cachedInputTokens": 4976640,
        "cacheWriteInputTokens": 0,
        "outputTokens": 6464,
        "reasoningOutputTokens": 1870
      },
      "max_observed_request_input_tokens": 65481,
      "elapsed_s": 1868.4370000000345,
      "turns_without_individual_usage": 0,
      "interrupted_turns_without_individual_usage": 0,
      "standard_api_equivalent_usd": 0.73252
    },
    {
      "run": "codex-shift-v1-20260930-r2",
      "primary": false,
      "outcome": "ERROR",
      "tokens": {
        "totalTokens": 4883995,
        "inputTokens": 4877485,
        "cachedInputTokens": 4758400,
        "cacheWriteInputTokens": 0,
        "outputTokens": 6510,
        "reasoningOutputTokens": 2153
      },
      "max_observed_request_input_tokens": 75362,
      "elapsed_s": 1868.9679999999935,
      "turns_without_individual_usage": 0,
      "interrupted_turns_without_individual_usage": 0,
      "standard_api_equivalent_usd": 0.77911
    }
  ],
  "claude_attempts": [
    {
      "run": "claude-shift-v1-20260930-r0",
      "primary": true,
      "outcome": "FAIL",
      "tokens_cli_reported": {
        "inputTokens": 186,
        "cacheCreationInputTokens": 53693,
        "cacheReadInputTokens": 3161503,
        "outputTokens": 7958,
        "thinkingTokens": 1208
      },
      "elapsed_s": 1865.015000000014,
      "cli_list_price_estimate_usd": 1.2217485999999997,
      "billed_usd": null
    },
    {
      "run": "claude-shift-v1-20260930-r1",
      "primary": true,
      "outcome": "FAIL",
      "tokens_cli_reported": {
        "inputTokens": 190,
        "cacheCreationInputTokens": 54865,
        "cacheReadInputTokens": 3245497,
        "outputTokens": 8666,
        "thinkingTokens": 1710
      },
      "elapsed_s": 1865.25,
      "cli_list_price_estimate_usd": 1.2620994000000005,
      "billed_usd": null
    },
    {
      "run": "claude-shift-v1-20260930-r2",
      "primary": true,
      "outcome": "FAIL",
      "tokens_cli_reported": {
        "inputTokens": 188,
        "cacheCreationInputTokens": 53686,
        "cacheReadInputTokens": 3197520,
        "outputTokens": 8537,
        "thinkingTokens": 1582
      },
      "elapsed_s": 1864.7029999999795,
      "cli_list_price_estimate_usd": 1.240484,
      "billed_usd": null
    }
  ],
  "claude_usage_conditions": "Signed-in Claude subscription; billed cost unavailable. Report the final native modelUsage cumulative totals, not the frozen adapter's partial output snapshots or sums of repeated results. CLI costUSD is Claude Code's own list-price estimate, not an invoice or a comparable metered API charge. No separate API-rate recalculation was made. Development pilot excluded.",
  "limits": "Retrospective cost analysis. Only the original seven metered API configurations are ranked. Codex and Claude Code actual costs are unavailable; neither native account is cost-ranked. The conditional Codex token-pricing scenario and Claude CLI list-price estimate do not establish a cost tie or advantage. Same workload, different models and timing conditions; later builds and concurrency differ from the original study. Superseded Codex wave and original error episodes are retained as overhead. Development pilots and this coding chat's work are excluded for all systems."
}
