{
  "agreementMetrics": null,
  "annotation": {
    "annotationTypes": [
      "classification",
      "span"
    ],
    "disagreementMetrics": "Available for supplied annotation sets; no multi-annotator study is claimed.",
    "metricsSupported": [
      "classificationAccuracy",
      "classificationF1",
      "classificationPrecision",
      "classificationRecall",
      "spanExactF1",
      "spanF1",
      "spanOverlapIoU",
      "spanPrecision",
      "spanRecall"
    ],
    "modality": "text",
    "modelCampaign": {
      "campaignId": "fc-e4d7e6adab0f",
      "coverage": {
        "cellsPending": 0,
        "cellsRun": 3,
        "complete": true,
        "completedAttempts": 30,
        "examsNotStarted": [],
        "missingAttempts": 0,
        "note": "Cells run against the cells this campaign declared it would run.",
        "pending": [],
        "planDeclared": true,
        "plannedAttempts": 30,
        "plannedCells": 3,
        "plannedExams": [
          "vvdex.annotation.text-labeling-1",
          "vvdex.annotation.structured-data-1",
          "vvdex.annotation.image-object-1"
        ],
        "plannedModels": [
          "cli/opencode-muse-spark"
        ],
        "rolloutsPerCell": 10
      },
      "examFingerprint": "e1abe51398b2d950a3415678aaf8ec796420dfb05ea89b3b64afd33c1e345fd3",
      "models": [
        {
          "attempts": 10,
          "delivery": "canonical text tools",
          "failed": 0,
          "invalid": 0,
          "laneErrors": 0,
          "model": "cli/opencode-muse-spark",
          "passed": 10,
          "quality": {
            "axes": [
              {
                "id": "classificationPrecision",
                "mean": 1.0,
                "n": 10
              },
              {
                "id": "classificationRecall",
                "mean": 1.0,
                "n": 10
              },
              {
                "id": "classificationF1",
                "mean": 1.0,
                "n": 10
              },
              {
                "id": "classificationAccuracy",
                "mean": 1.0,
                "n": 10
              },
              {
                "id": "spanPrecision",
                "mean": 1.0,
                "n": 10
              },
              {
                "id": "spanRecall",
                "mean": 1.0,
                "n": 10
              },
              {
                "id": "spanF1",
                "mean": 1.0,
                "n": 10
              },
              {
                "id": "spanOverlapIoU",
                "mean": 1.0,
                "n": 10
              },
              {
                "id": "spanExactF1",
                "mean": 1.0,
                "n": 10
              }
            ],
            "measuredAttempts": 10,
            "scoreMean": 1.0
          },
          "stages": {
            "edit": 10,
            "graded": 10,
            "hidden_pass": 10,
            "inspect": 10,
            "submit": 10,
            "test": 10
          },
          "withheld": 0
        }
      ],
      "reportUrl": "/reports/fc-e4d7e6adab0f/",
      "summarySha256": "7a5d1a00820ffb56bc51c8728e14c9b9d5e5348c962d902e6b673a63873f6b15",
      "summaryUrl": "/reports/fc-e4d7e6adab0f/summary.json"
    },
    "purpose": "An invented message is reviewed for classification and entity spans with consistent offsets.",
    "rubricVersion": "1.0"
  },
  "assetSource": null,
  "certificationReceiptSha256": "cd09865ac07c85cf1bb3c6eacee46d90754c915788b6101248a284010ddb1d67",
  "certificationState": "certified",
  "disclosure": "Aggregate certification evidence only; reference annotations and reusable evaluation material withheld.",
  "errorClasses": [
    "malformed input",
    "schema mismatch",
    "out of bounds",
    "missing annotation",
    "temporal conflict",
    "inconsistent tracking",
    "reference mismatch"
  ],
  "examFingerprint": "e1abe51398b2d950a3415678aaf8ec796420dfb05ea89b3b64afd33c1e345fd3",
  "examId": "vvdex.annotation.text-labeling-1",
  "limitations": "Small constructed engineering fixtures, not a production dataset or professional annotation engagement. Video uses sparse frame boxes, not dense tracking. Audio contains tones and silence, not speech or speakers. CVAT XML 1.1 and Label Studio JSON cover tested box subsets only; unsupported or lossy conversions are refused. Polygon comparison requires the implemented vertex representation. The attached model campaign measures these fixed tasks only; general capability and annotator population performance are not inferred.",
  "modelCampaign": {
    "campaignId": "fc-e4d7e6adab0f",
    "coverage": {
      "cellsPending": 0,
      "cellsRun": 3,
      "complete": true,
      "completedAttempts": 30,
      "examsNotStarted": [],
      "missingAttempts": 0,
      "note": "Cells run against the cells this campaign declared it would run.",
      "pending": [],
      "planDeclared": true,
      "plannedAttempts": 30,
      "plannedCells": 3,
      "plannedExams": [
        "vvdex.annotation.text-labeling-1",
        "vvdex.annotation.structured-data-1",
        "vvdex.annotation.image-object-1"
      ],
      "plannedModels": [
        "cli/opencode-muse-spark"
      ],
      "rolloutsPerCell": 10
    },
    "examFingerprint": "e1abe51398b2d950a3415678aaf8ec796420dfb05ea89b3b64afd33c1e345fd3",
    "models": [
      {
        "attempts": 10,
        "delivery": "canonical text tools",
        "failed": 0,
        "invalid": 0,
        "laneErrors": 0,
        "model": "cli/opencode-muse-spark",
        "passed": 10,
        "quality": {
          "axes": [
            {
              "id": "classificationPrecision",
              "mean": 1.0,
              "n": 10
            },
            {
              "id": "classificationRecall",
              "mean": 1.0,
              "n": 10
            },
            {
              "id": "classificationF1",
              "mean": 1.0,
              "n": 10
            },
            {
              "id": "classificationAccuracy",
              "mean": 1.0,
              "n": 10
            },
            {
              "id": "spanPrecision",
              "mean": 1.0,
              "n": 10
            },
            {
              "id": "spanRecall",
              "mean": 1.0,
              "n": 10
            },
            {
              "id": "spanF1",
              "mean": 1.0,
              "n": 10
            },
            {
              "id": "spanOverlapIoU",
              "mean": 1.0,
              "n": 10
            },
            {
              "id": "spanExactF1",
              "mean": 1.0,
              "n": 10
            }
          ],
          "measuredAttempts": 10,
          "scoreMean": 1.0
        },
        "stages": {
          "edit": 10,
          "graded": 10,
          "hidden_pass": 10,
          "inspect": 10,
          "submit": 10,
          "test": 10
        },
        "withheld": 0
      }
    ],
    "reportUrl": "/reports/fc-e4d7e6adab0f/",
    "summarySha256": "7a5d1a00820ffb56bc51c8728e14c9b9d5e5348c962d902e6b673a63873f6b15",
    "summaryUrl": "/reports/fc-e4d7e6adab0f/summary.json"
  },
  "referenceMetrics": {
    "classificationAccuracy": 1.0,
    "classificationF1": 1.0,
    "classificationPrecision": 1.0,
    "classificationRecall": 1.0,
    "spanExactF1": 1.0,
    "spanF1": 1.0,
    "spanOverlapIoU": 1.0,
    "spanPrecision": 1.0,
    "spanRecall": 1.0
  },
  "schema": "vvdex.forge.annotation-certification-report/v1"
}
