{
  "schema": "konstant.website.consumer-capability-evidence/v1",
  "consumer": "Konstant",
  "provider": "Precise",
  "consumer_objective": "Improve how Gary carries source-backed information into answers; choose which evidence representation to keep testing.",
  "provider_method": "Exact finite-game Shapley attribution and pair interactions",
  "integration": "Existing Konstant extraction evaluation adapter calls the actual Precise TypeScript/Python implementation in a separate local source checkout.",
  "measured_at": "2026-07-17T14:29:31.825Z",
  "replayed_at": "2026-09-06T21:59:33.790Z",
  "status": "Historical internal product evaluation; exact attribution replayed locally.",
  "rows": [
    {
      "key": "",
      "label": "Baseline",
      "exact": 78,
      "total": 178,
      "partial": 80,
      "missed": 20,
      "coverage": 0.43820224719101125,
      "quality_credit": 0.6629213483146067,
      "model_visible_bytes": 562590
    },
    {
      "key": "quality_filter",
      "label": "Filter only",
      "exact": 84,
      "total": 178,
      "partial": 76,
      "missed": 18,
      "coverage": 0.47191011235955055,
      "quality_credit": 0.6853932584269663,
      "model_visible_bytes": 545528
    },
    {
      "key": "lossless_ledger",
      "label": "Compact format only",
      "exact": 95,
      "total": 178,
      "partial": 58,
      "missed": 25,
      "coverage": 0.5337078651685393,
      "quality_credit": 0.6966292134831461,
      "model_visible_bytes": 354091
    },
    {
      "key": "quality_filter|lossless_ledger",
      "label": "Both changes",
      "exact": 80,
      "total": 178,
      "partial": 72,
      "missed": 26,
      "coverage": 0.449438202247191,
      "quality_credit": 0.651685393258427,
      "model_visible_bytes": 339086
    }
  ],
  "shapley": {
    "quality_filter": -0.0252808988764045,
    "lossless_ledger": 0.03651685393258425
  },
  "interaction": -0.11797752808988754,
  "implementation_sha256": "sha256:b8739eb870727f12bde8394715f22e493775b0a0904a56ad82961a0161d98cf2",
  "decision": {
    "status": "Recorded development decision with a completed follow-up replication",
    "text": "Retain the compact format without this filter as the candidate for further testing; stop combining the two in this experiment.",
    "production_promotion": false,
    "followup_receipt_sha256": "sha256:dbe986bc2641fd2bc23896082aced17387c890e3340be24612358b1a0da5d10b"
  },
  "ownership": {
    "konstant": "Source records, answer-quality criteria, measured outcomes, integration and product decisions.",
    "precise": "Attribution mathematics and its implementation.",
    "boombox": "Not invoked in this attribution replay; separate local operating demonstration available.",
    "arranger": "Not invoked in this attribution replay; separate local next-experiment study available."
  },
  "verification": {
    "candidate_receipt_hashes": true,
    "all_four_cells_bound_to_source_receipts": true,
    "all_displayed_counts_match_evaluation_aggregates": true,
    "actual_provider_replay": true,
    "historical_attribution_match": true,
    "parent_run_manifest": "The parent run manifest was regenerated after the historical experiment. Its current digest differs from the historical reference. The unchanged individual candidate receipts still match both the projection receipts and every game cell; verification uses those exact records."
  },
  "limits": [
    "Two source documents and 178 evaluation obligations in one development comparison.",
    "Exact coverage is a recorded evaluation judgment, not a percentage of universally correct answers.",
    "The September replay recalculates attribution over July measurements; it does not regenerate or independently rejudge the answers.",
    "No customer economic outcome, production promotion, hosted cross-company authentication, payment or live Gary session is established.",
    "Only aggregate product-evaluation results and opaque evidence identities are exported. Source records, document names and individual judgments remain private."
  ],
  "files": [
    {
      "file": "coverage-input.json",
      "sha256": "0f9da1a15dc9e151ad8dfb29374e98ae3cd77e9fa39509f41dee6264f9dff952"
    },
    {
      "file": "coverage-result.json",
      "sha256": "7b95abf775c25e65287f4f38c76bb4a99f18fb00482a92d6cb9a6f2ee192a1cb"
    },
    {
      "file": "quality-input.json",
      "sha256": "af27457506697ee05a8d3ad125224dc8a6b4ec571a2b058a9d9ccefebecd01fa"
    },
    {
      "file": "quality-result.json",
      "sha256": "ff955f683201d8532392724a5fd4e9cde4797dca1a49458f6ace3dd7e76668e8"
    },
    {
      "file": "bytes-input.json",
      "sha256": "2b22830a3ddc2214ab8f8bd83f1623567f8b1de3b8150b59c817a86da1347e7b"
    },
    {
      "file": "bytes-result.json",
      "sha256": "45beb08b8d90d09b0ecc3f4b562f5342d9d68db0f08117417b1efc450b8c4a21"
    },
    {
      "file": "arranger-study.json",
      "sha256": "123942bf847299570a30a7cb64334c2b06d4dc56665308e7a55849429737997a"
    }
  ],
  "arranger_study": {
    "status": "Executed local integration with illustrative probabilities; research study",
    "recommended_experiment": "joint-batch",
    "placebo_stops": true,
    "production_changed": false,
    "suite_hash": "sha256:c85000202a3706ec6e5c6a4dbde1165779c92c8a7f912e63d629cc41a7e19b01"
  }
}
