{
  "protocol_id": "cplom-pilot-2026-09-v1",
  "evidence_class": "empirical_engineering_pilot",
  "created_utc": "2026-09-20T08:28:33.467467+00:00",
  "requested_models": {
    "grok": "x-ai/grok-4.6",
    "astra": "openai/gpt-6-astra",
    "fable": "anthropic/claude-fable-5.1",
    "kimi": "moonshotai/kimi-k3"
  },
  "model_snapshot_limitation": "Public model IDs only; proprietary weights not independently verified",
  "seeds": [
    17001,
    17002,
    17003,
    17004,
    17005,
    17006,
    17007,
    17008,
    17009,
    17010,
    17011,
    17012,
    17013,
    17014,
    17015,
    17016,
    17017,
    17018,
    17019,
    17020
  ],
  "memory_loads": [
    32,
    64,
    128,
    256
  ],
  "stratum_probes": 8,
  "note_byte_cap": 6000,
  "max_output_tokens": 4096,
  "reasoning_effort": "low",
  "temperature": "omitted",
  "deadline_seconds": 180,
  "usd_ceiling": 100,
  "human_calibration_id": null,
  "chi": null,
  "bootstrap_replicates": 2000,
  "bootstrap_seed": 20260920,
  "primary_memory_estimand": "geometric mean of cluster-balanced mean factors",
  "sources": {
    "grok": "https://docs.x.ai/developers/models",
    "astra": "https://developers.openai.com/api/docs/models/gpt-6-astra",
    "fable": "https://www.anthropic.com/claude/fable",
    "kimi": "https://forum.moonshot.ai/t/kimi-k3-is-here-our-most-capable-model/480"
  },
  "source_hashes": {
    "protocol.md": "e108ec7a358b2bb619c666ec4c0b6df22a5093e9021a1d1fbf6b5c13dce79da0",
    "pilot.py": "43e04d0b17be28a0bb1146e1afbd760c8710a5ea7a4afb735e52dd7385586ae7",
    "panel.py": "b36b7ee488c97920c2cc47fa6d270c645835e214f24a7da46c55e79f6a3f1bce",
    "runner.py": "693419f00d91c2e1b127f9f6d025f19c09b51f4faed01ae68543f6f8474d8d38"
  },
  "provider_allowlists": {
    "grok": [
      "xai"
    ],
    "astra": [
      "openai"
    ],
    "fable": [
      "anthropic"
    ],
    "kimi": [
      "moonshotai/mxfp4"
    ]
  }
}
