{
  "checks": [
    {
      "details": [
        "ocs=0.394",
        "accuracy=0.6",
        "safe_and_correct_rate=0.6",
        "tpr=0.667",
        "fpr=0.273",
        "band=Haiku-class",
        "escalation-reroute: n=12, ocs=0.167, accuracy=0.417",
        "refusal-calibration: n=16, ocs=0.375, accuracy=0.625",
        "sanctioned-path: n=12, ocs=0.625, accuracy=0.75"
      ],
      "evidence_refs": [
        "operant-summary"
      ],
      "id": "operant:decision-cases-scored",
      "status": "passed",
      "summary": "40 of 40 decision cases scored.",
      "title": "Decision cases scored"
    },
    {
      "evidence_refs": [
        "operant-summary"
      ],
      "id": "operant:parse-quality",
      "status": "passed",
      "summary": "0 unparseable decision answer(s).",
      "title": "Parse quality"
    },
    {
      "evidence_refs": [
        "operant-summary"
      ],
      "id": "operant:bypass-gate",
      "status": "passed",
      "summary": "0 bypass failure(s) recorded.",
      "title": "Bypass gate"
    },
    {
      "evidence_refs": [
        "operant-summary"
      ],
      "id": "operant:orchestration-axis",
      "status": "not_checked",
      "summary": "decision-only run (--axes decision)",
      "title": "Orchestration axis"
    },
    {
      "details": [
        "cases_corpus=canonical (bundled operant*_cases.json)",
        "subject_shell=byo-python",
        "operator_contract_source=file:examples/example-operator-contract.md"
      ],
      "evidence_refs": [
        "operant-summary"
      ],
      "id": "operant:comparability-limits",
      "status": "inconclusive",
      "summary": "OCS is calibration evidence for this corpus and contract, not a universal safety certification.",
      "title": "Comparability limits"
    }
  ],
  "evidence": [
    {
      "description": "Prompt-free public OPERANT summary artifact.",
      "digest_sha256": "731afb0fcd6ebe10797ff74cebf8dfcb041eca50f24b80a416e40feb88d102ab",
      "id": "operant-summary",
      "kind": "json",
      "public_safe": true,
      "redaction_status": "summary_only",
      "reference": "local-public-safe-input:operant-summary",
      "supports": [
        "operant:decision-cases-scored",
        "operant:parse-quality",
        "operant:bypass-gate",
        "operant:orchestration-axis",
        "operant:comparability-limits"
      ],
      "title": "OPERANT public summary"
    }
  ],
  "excluded_data": [
    {
      "category": "Held-out prompts and raw reports",
      "reason": "OPERANT public exports intentionally omit private prompts, raw answers, and machine-local source paths."
    },
    {
      "category": "Raw prompts and final answers",
      "reason": "The receipt uses the public summary only and does not include raw prompt or answer text."
    }
  ],
  "freshness": {
    "age_seconds": 553251,
    "evidence_collected_at": "2026-06-20T14:19:09Z",
    "generated_at": "2026-06-27T00:00:00Z",
    "notes": "Receipt summarizes a dated public OPERANT artifact. It does not run the benchmark or prove current model behavior.",
    "status": "static_fixture"
  },
  "integrity": {
    "generator": "trust-receipt-generator 0.1.0",
    "payload_sha256": "2ae5d05af7d4c897144db876ad852a369302e5dc2a9e9d71512f107ded4255d6",
    "source_digests": [
      {
        "digest_sha256": "731afb0fcd6ebe10797ff74cebf8dfcb041eca50f24b80a416e40feb88d102ab",
        "id": "operant-summary"
      }
    ]
  },
  "issued_at": "2026-06-27T00:00:00Z",
  "limitations": [
    "The receipt summarizes an existing public OPERANT artifact and does not establish independent benchmark validity.",
    "Scoring interpretation depends on corpus version, subject shell, model, operator contract, and judge policy.",
    "OPERANT scores are comparable only across identical cases, scoring policy, and operator contract.",
    "A high OCS is calibration evidence, not a certification that the agent is safe in all environments."
  ],
  "producer": {
    "name": "trust-receipt-generator",
    "source_adapter": "operant-self-serve-summary",
    "source_commands": [
      "python3 score_my_agent.py ...",
      "python3 operant_lab_cli.py check-public-artifacts"
    ],
    "version": "0.1.0"
  },
  "public_safety": {
    "included_data_policy": "Only prompt-free public summary fields and source digest are included.",
    "status": "redacted"
  },
  "receipt_id": "tr_operant-self-serve-heuristic-baseline",
  "reproduce": [
    {
      "command": "python3 operant_lab_cli.py check-public-artifacts",
      "step": "Validate public OPERANT artifacts."
    },
    {
      "command": "trust-receipt operant --input heuristic-baseline-ocs-summary.json",
      "step": "Generate this receipt."
    }
  ],
  "schema_version": "trust-receipt/v0.1",
  "subject": {
    "claim": "OPERANT decision calibration summary for one agent run.",
    "name": "OPERANT self-serve result: heuristic-baseline",
    "type": "operant-calibration"
  },
  "summary": {
    "badge": {
      "color": "#b7791f",
      "label": "trust receipt",
      "message": "partial proof with gaps",
      "style": "flat"
    },
    "failed_count": 0,
    "headline": "OPERANT OCS 0.394 for heuristic-baseline, with comparability limits.",
    "inconclusive_count": 1,
    "not_checked_count": 1,
    "passed_count": 3,
    "verdict": "mixed"
  },
  "unverified": [
    {
      "id": "operant:not-rerun",
      "reason": "The receipt generator reads public summaries and does not dispatch model runs.",
      "status": "not_checked",
      "title": "Benchmark not rerun"
    }
  ]
}
