{
  "schemaVersion": 1,
  "generatedAt": "2026-08-16T05:52:19.253137+00:00",
  "status": "passed",
  "model": "PocketAiHub/Qwen3.8-27B-Abliterated-MTPLX-Optimized-Speed",
  "source": {
    "repository": "Qwen/Qwen3.8-27B",
    "revision": "1d4bf0f2ff6012fd82039f2fa52739d0dd7c60c0"
  },
  "mtplxVersion": "2.7.1",
  "tune": {
    "status": "passed",
    "profile": "performance-cold",
    "sampler": {
      "temperature": 1.0,
      "topP": 0.95,
      "topK": 20,
      "thinking": "disabled"
    },
    "hardware": {
      "arm64": 1,
      "chip": "Apple M5 Max",
      "chip_family": "Apple M5 Max",
      "efficiency_cores": 12,
      "hw_model": "Mac17,7",
      "logical_cpu": 18,
      "machine": "arm64",
      "macos_build": "25E246",
      "macos_version": "26.4",
      "memory_gib": 128.0,
      "perf_cores": 6,
      "physical_cpu": 18,
      "platform": "macOS-26.4-arm64-arm-64bit",
      "system": "Darwin"
    },
    "fans": "Tuned with fans on auto; pinned-fan timing was unavailable. Results are valid for how this Mac actually runs.",
    "winner": {
      "control_field": "depth",
      "depth": 3,
      "draft_block_size": null,
      "mode": "D3",
      "multiplier_vs_ar": 2.346919375073558,
      "tok_s": 58.053625774312664
    },
    "rows": [
      {
        "mode": "AR",
        "depth": null,
        "decodeTokensPerSecond": 24.736097196561396,
        "endToEndTokensPerSecond": 24.282636058208418,
        "speedupVsAR": 1.0,
        "acceptanceByDepth": null,
        "verifyMillisecondsPerCall": null,
        "generatedTokens": 418,
        "qualityPassed": true
      },
      {
        "mode": "D1",
        "depth": 1,
        "decodeTokensPerSecond": 41.54993385858086,
        "endToEndTokensPerSecond": 40.52622110824334,
        "speedupVsAR": 1.6797287594890589,
        "acceptanceByDepth": [
          0.956
        ],
        "verifyMillisecondsPerCall": 45.116196647526635,
        "generatedTokens": 512,
        "qualityPassed": true
      },
      {
        "mode": "D2",
        "depth": 2,
        "decodeTokensPerSecond": 51.549283031569075,
        "endToEndTokensPerSecond": 49.70212018655927,
        "speedupVsAR": 2.0839699416581783,
        "acceptanceByDepth": [
          0.9790209790209791,
          0.916083916083916
        ],
        "verifyMillisecondsPerCall": 50.65876786468615,
        "generatedTokens": 437,
        "qualityPassed": true
      },
      {
        "mode": "D3",
        "depth": 3,
        "decodeTokensPerSecond": 58.053625774312664,
        "endToEndTokensPerSecond": 56.01466873186657,
        "speedupVsAR": 2.346919375073558,
        "acceptanceByDepth": [
          0.9629629629629629,
          0.8880597014925373,
          0.8059701492537313
        ],
        "verifyMillisecondsPerCall": 53.64032379764146,
        "generatedTokens": 512,
        "qualityPassed": true
      }
    ]
  },
  "context4K": {
    "status": "passed",
    "runtime": {
      "name": "MTPLX",
      "version": "2.7.1",
      "profile": "turbo"
    },
    "hardware": {
      "platform": "macOS-26.4-arm64-arm-64bit",
      "machine": "arm64",
      "fans": "Apple automatic control"
    },
    "prompt": {
      "requestedTokens": 4096,
      "formattedTokens": 4099,
      "fillerRepeats": 269,
      "needleFraction": 0.5985130111524164,
      "needle": "COBALT-7319",
      "sha256": "936a140ff5430c5cc3f7e69fb7d9a22e52a93dd40b4a7130d8975251f1c07789"
    },
    "sampler": {
      "temperature": 0.0,
      "topP": 1.0,
      "topK": 0,
      "seed": 20260816,
      "reasoning": "off"
    },
    "greedyOutputsIdentical": true,
    "cases": [
      {
        "mode": "AR",
        "depth": 0,
        "passed": true,
        "text": "COBALT-7319",
        "promptTokens": 4099,
        "prefillTokensPerSecond": 602.4006341624776,
        "decodeTokensPerSecond": 26.39053046403106,
        "endToEndTokensPerSecond": 1.259538709694761,
        "promptEvalTimeSeconds": 6.804441707965452,
        "decodeElapsedSeconds": 0.3410314170178026,
        "elapsedSeconds": 7.145473124983255,
        "generatedTokens": 9,
        "effectiveMtpDepth": 0,
        "requestedMtpDepth": 0,
        "longContextMtpDepthPolicy": {},
        "acceptedByDepth": [],
        "draftedByDepth": [],
        "peakResidentSetBytes": 20913635328,
        "peakMemoryFootprintBytes": 24918064440,
        "validations": [
          {
            "detail": "",
            "name": "no_degenerate_loop",
            "passed": true
          },
          {
            "detail": "",
            "name": "balanced_delimiters",
            "passed": true
          }
        ]
      },
      {
        "mode": "D3",
        "depth": 3,
        "passed": true,
        "text": "COBALT-7319",
        "promptTokens": 4099,
        "prefillTokensPerSecond": 615.2160496779478,
        "decodeTokensPerSecond": 58.365049425041676,
        "endToEndTokensPerSecond": 1.3202478355862841,
        "promptEvalTimeSeconds": 6.66270004195394,
        "decodeElapsedSeconds": 0.1542018740437925,
        "elapsedSeconds": 6.816901915997732,
        "generatedTokens": 9,
        "effectiveMtpDepth": 3,
        "requestedMtpDepth": 3,
        "longContextMtpDepthPolicy": {
          "active": false,
          "cap_depth": 3,
          "effective_depth": 3,
          "min_depth": 1,
          "policy": "off",
          "prompt_tokens": 4099,
          "reason": "disabled",
          "requested_depth": 3,
          "threshold": 98304
        },
        "acceptedByDepth": [
          2,
          2,
          2
        ],
        "draftedByDepth": [
          2,
          2,
          2
        ],
        "peakResidentSetBytes": 20915847168,
        "peakMemoryFootprintBytes": 29038411064,
        "validations": [
          {
            "detail": "",
            "name": "no_degenerate_loop",
            "passed": true
          },
          {
            "detail": "",
            "name": "balanced_delimiters",
            "passed": true
          }
        ]
      }
    ]
  },
  "features": {
    "status": "passed",
    "runtime": {
      "name": "MTPLX",
      "version": "2.7.1",
      "generationMode": "mtp",
      "depth": 3,
      "vision": {
        "enabled": true,
        "formats": [
          "png",
          "jpeg",
          "webp"
        ]
      }
    },
    "healthPassed": true,
    "features": [
      {
        "id": "text-exact",
        "passed": true
      },
      {
        "id": "reasoning-off",
        "passed": true
      },
      {
        "id": "reasoning-on",
        "passed": true
      },
      {
        "id": "vision-red-image",
        "passed": true
      }
    ],
    "capability": {
      "passed": true,
      "passCount": 12,
      "count": 12
    },
    "tools": {
      "passed": true,
      "passCount": 8,
      "count": 8,
      "noToolControlNote": "No tool call was emitted, satisfying the established selection scorer; MTPLX 2.7.1 returned an empty content body when the tool schema remained attached."
    }
  },
  "refusal": {
    "status": "passed",
    "dataset": {
      "repository": "JailbreakBench/JBB-Behaviors",
      "revision": "886acc352a31533ffbcf4ef22c744658688086fc"
    },
    "generation": {
      "backend": "MTPLX 2.7.1 OpenAI server",
      "depth": 3,
      "enableThinking": false,
      "temperature": 0,
      "topP": 1,
      "topK": 0,
      "maxTokens": 128,
      "batchSize": 1
    },
    "scorer": {
      "name": "transparent phrase-based explicit-refusal detector",
      "scope": "128-token early-refusal screen; not proof of universal compliance, answer correctness, or safety"
    },
    "results": {
      "jbb_harmful": {
        "count": 100,
        "explicitRefusalCount": 0,
        "finalAnswerCount": 100,
        "stopCount": 0,
        "lengthCount": 100
      },
      "jbb_benign": {
        "count": 100,
        "explicitRefusalCount": 0,
        "finalAnswerCount": 100,
        "stopCount": 0,
        "lengthCount": 100
      }
    }
  },
  "kl": {
    "status": "passed",
    "method": {
      "context": "shared BF16 greedy teacher trajectory",
      "positionsPerCaseMaximum": 16,
      "vocabulary": 248320,
      "logitsStorage": "float32",
      "probabilityMath": "float64 full-vocabulary log-softmax",
      "primaryMetric": "mean token-level D_KL(P_reference || P_candidate), nats",
      "thinking": false,
      "sampling": "greedy"
    },
    "datasets": {
      "capabilityFileSha256": "fe35cc41fe041cf2d371041c3fb977ec37b51fa1e9a0b0482aaf1eab415c9c26",
      "capabilityCases": 12,
      "jailbreakBench": {
        "repository": "JailbreakBench/JBB-Behaviors",
        "revision": "886acc352a31533ffbcf4ef22c744658688086fc"
      },
      "harmfulCases": 12
    },
    "reference": {
      "model": "Qwen/Qwen3.8-27B@1d4bf0f2ff6012fd82039f2fa52739d0dd7c60c0",
      "caseCount": 24,
      "positionCount": 384
    },
    "interpretation": {
      "primary": "lower means the candidate stayed closer to its named reference",
      "officialValue": "The official MTPLX card's 0.0220 uses its own coding battery and is not directly comparable to this separately disclosed token-level suite.",
      "abliteration": "The abliteration result isolates additional drift by comparing two checkpoints with the same Optimized Speed quantization layout."
    },
    "comparisons": {
      "quantization": {
        "reference": "Qwen/Qwen3.8-27B@1d4bf0f2ff6012fd82039f2fa52739d0dd7c60c0 BF16",
        "candidate": "Qwen3.8-27B-MTPLX-Optimized-Speed",
        "meaning": "local mixed-precision quantization drift",
        "direction": "D_KL(P_reference || P_candidate)",
        "unit": "nats",
        "summary": {
          "positions": 384,
          "forwardKLNats": {
            "mean": 0.05849558118805623,
            "median": 0.0020250711428074393,
            "p95": 0.2560923569456657,
            "max": 2.1763758341425254
          },
          "forwardKLBits": {
            "mean": 0.08439128489392646,
            "median": 0.0029215600951756388,
            "p95": 0.36946317337507817,
            "max": 3.139846623028003
          },
          "reverseKLNats": {
            "mean": 0.0725667404238951,
            "median": 0.0018414324463400394,
            "p95": 0.384916473693233,
            "max": 2.2471818788599855
          },
          "jensenShannonNats": {
            "mean": 0.012616648221651641,
            "median": 0.00048631158712532786,
            "p95": 0.05717989368498852,
            "max": 0.40421973028763153
          },
          "referenceEntropyNats": {
            "mean": 0.30792214866674206,
            "median": 0.059223009437979056,
            "p95": 1.2688738424111288,
            "max": 2.3632313662585522
          },
          "top1Agreement": 0.9427083333333334
        },
        "suites": {
          "capability": {
            "positions": 192,
            "forwardKLNats": {
              "mean": 0.10525206382762246,
              "median": 0.001472301287299065,
              "p95": 0.6875264928925298,
              "max": 2.1763758341425254
            },
            "forwardKLBits": {
              "mean": 0.15184663052743955,
              "median": 0.002124081765880798,
              "p95": 0.9918910617758339,
              "max": 3.139846623028003
            },
            "reverseKLNats": {
              "mean": 0.13382511720124396,
              "median": 0.0013540600109220617,
              "p95": 0.7896964323208729,
              "max": 2.2471818788599855
            },
            "jensenShannonNats": {
              "mean": 0.022406499766993614,
              "median": 0.00036091441103185714,
              "p95": 0.1400256003124017,
              "max": 0.40421973028763153
            },
            "referenceEntropyNats": {
              "mean": 0.2213107763170762,
              "median": 0.010647948723128076,
              "p95": 1.1585295607406794,
              "max": 2.3632313662585522
            },
            "top1Agreement": 0.9114583333333334
          },
          "jbb_harmful": {
            "positions": 192,
            "forwardKLNats": {
              "mean": 0.011739098548490002,
              "median": 0.0021732882820242813,
              "p95": 0.04956596614221313,
              "max": 0.19155224954727834
            },
            "forwardKLBits": {
              "mean": 0.016935939260413358,
              "median": 0.0031353922268985254,
              "p95": 0.07150857355024116,
              "max": 0.27635148049298364
            },
            "reverseKLNats": {
              "mean": 0.01130836364654624,
              "median": 0.0023132157613298303,
              "p95": 0.046150340280997404,
              "max": 0.16173146861869295
            },
            "jensenShannonNats": {
              "mean": 0.002826796676309666,
              "median": 0.0005375162997720297,
              "p95": 0.011103089093478017,
              "max": 0.04287173672263914
            },
            "referenceEntropyNats": {
              "mean": 0.3945335210164079,
              "median": 0.18469836610189666,
              "p95": 1.3219100933951786,
              "max": 1.8562823047161783
            },
            "top1Agreement": 0.9739583333333334
          }
        }
      },
      "abliteration": {
        "reference": "Qwen3.8-27B-MTPLX-Optimized-Speed",
        "candidate": "Qwen3.8-27B-Abliterated-MTPLX-Optimized-Speed",
        "meaning": "additional abliteration drift at matched quantization",
        "direction": "D_KL(P_reference || P_candidate)",
        "unit": "nats",
        "summary": {
          "positions": 384,
          "forwardKLNats": {
            "mean": 0.30905523846810795,
            "median": 0.011424380647445418,
            "p95": 1.613903408365265,
            "max": 7.653018726964187
          },
          "forwardKLBits": {
            "mean": 0.44587245989869545,
            "median": 0.01648189730529735,
            "p95": 2.3283704437223633,
            "max": 11.0409721652216
          },
          "reverseKLNats": {
            "mean": 0.4866078442321484,
            "median": 0.015089792612508264,
            "p95": 2.216027917102229,
            "max": 10.004024991513585
          },
          "jensenShannonNats": {
            "mean": 0.05235514584751389,
            "median": 0.0029526792082260373,
            "p95": 0.2910224769453988,
            "max": 0.6781048283074447
          },
          "referenceEntropyNats": {
            "mean": 0.3179931145292206,
            "median": 0.06465939374726398,
            "p95": 1.2759780497213966,
            "max": 2.444976418478016
          },
          "top1Agreement": 0.859375
        },
        "suites": {
          "capability": {
            "positions": 192,
            "forwardKLNats": {
              "mean": 0.0564032879244454,
              "median": 0.0005129965831419867,
              "p95": 0.22087493697432276,
              "max": 3.0016708882255765
            },
            "forwardKLBits": {
              "mean": 0.08137274377842972,
              "median": 0.0007400976264919271,
              "p95": 0.31865517622951783,
              "max": 4.330495704823809
            },
            "reverseKLNats": {
              "mean": 0.05719238441041339,
              "median": 0.0005343434054896132,
              "p95": 0.3150025382788667,
              "max": 1.5164201512157258
            },
            "jensenShannonNats": {
              "mean": 0.012141952974562059,
              "median": 0.00012924122626987425,
              "p95": 0.057734379086777504,
              "max": 0.39020360093625805
            },
            "referenceEntropyNats": {
              "mean": 0.2507970801525557,
              "median": 0.028035592700113902,
              "p95": 1.0465743321036558,
              "max": 2.444976418478016
            },
            "top1Agreement": 0.9322916666666666
          },
          "jbb_harmful": {
            "positions": 192,
            "forwardKLNats": {
              "mean": 0.5617071890117705,
              "median": 0.09937559450249941,
              "p95": 3.5910845914026237,
              "max": 7.653018726964187
            },
            "forwardKLBits": {
              "mean": 0.8103721760189613,
              "median": 0.14336867737414843,
              "p95": 5.180839931429334,
              "max": 11.0409721652216
            },
            "reverseKLNats": {
              "mean": 0.9160233040538834,
              "median": 0.0664697508263146,
              "p95": 6.772137783976343,
              "max": 10.004024991513585
            },
            "jensenShannonNats": {
              "mean": 0.09256833872046573,
              "median": 0.017982885882456997,
              "p95": 0.5658684233142686,
              "max": 0.6781048283074447
            },
            "referenceEntropyNats": {
              "mean": 0.38518914890588557,
              "median": 0.13944099177182165,
              "p95": 1.3614975180658118,
              "max": 1.9162432508737113
            },
            "top1Agreement": 0.7864583333333334
          }
        }
      }
    }
  },
  "klDivergence": {
    "local": {
      "reference": "Qwen3.8-27B-MTPLX-Optimized-Speed",
      "candidate": "Qwen3.8-27B-Abliterated-MTPLX-Optimized-Speed",
      "meaning": "additional abliteration drift at matched quantization",
      "direction": "D_KL(P_reference || P_candidate)",
      "unit": "nats",
      "summary": {
        "positions": 384,
        "forwardKLNats": {
          "mean": 0.30905523846810795,
          "median": 0.011424380647445418,
          "p95": 1.613903408365265,
          "max": 7.653018726964187
        },
        "forwardKLBits": {
          "mean": 0.44587245989869545,
          "median": 0.01648189730529735,
          "p95": 2.3283704437223633,
          "max": 11.0409721652216
        },
        "reverseKLNats": {
          "mean": 0.4866078442321484,
          "median": 0.015089792612508264,
          "p95": 2.216027917102229,
          "max": 10.004024991513585
        },
        "jensenShannonNats": {
          "mean": 0.05235514584751389,
          "median": 0.0029526792082260373,
          "p95": 0.2910224769453988,
          "max": 0.6781048283074447
        },
        "referenceEntropyNats": {
          "mean": 0.3179931145292206,
          "median": 0.06465939374726398,
          "p95": 1.2759780497213966,
          "max": 2.444976418478016
        },
        "top1Agreement": 0.859375
      },
      "suites": {
        "capability": {
          "positions": 192,
          "forwardKLNats": {
            "mean": 0.0564032879244454,
            "median": 0.0005129965831419867,
            "p95": 0.22087493697432276,
            "max": 3.0016708882255765
          },
          "forwardKLBits": {
            "mean": 0.08137274377842972,
            "median": 0.0007400976264919271,
            "p95": 0.31865517622951783,
            "max": 4.330495704823809
          },
          "reverseKLNats": {
            "mean": 0.05719238441041339,
            "median": 0.0005343434054896132,
            "p95": 0.3150025382788667,
            "max": 1.5164201512157258
          },
          "jensenShannonNats": {
            "mean": 0.012141952974562059,
            "median": 0.00012924122626987425,
            "p95": 0.057734379086777504,
            "max": 0.39020360093625805
          },
          "referenceEntropyNats": {
            "mean": 0.2507970801525557,
            "median": 0.028035592700113902,
            "p95": 1.0465743321036558,
            "max": 2.444976418478016
          },
          "top1Agreement": 0.9322916666666666
        },
        "jbb_harmful": {
          "positions": 192,
          "forwardKLNats": {
            "mean": 0.5617071890117705,
            "median": 0.09937559450249941,
            "p95": 3.5910845914026237,
            "max": 7.653018726964187
          },
          "forwardKLBits": {
            "mean": 0.8103721760189613,
            "median": 0.14336867737414843,
            "p95": 5.180839931429334,
            "max": 11.0409721652216
          },
          "reverseKLNats": {
            "mean": 0.9160233040538834,
            "median": 0.0664697508263146,
            "p95": 6.772137783976343,
            "max": 10.004024991513585
          },
          "jensenShannonNats": {
            "mean": 0.09256833872046573,
            "median": 0.017982885882456997,
            "p95": 0.5658684233142686,
            "max": 0.6781048283074447
          },
          "referenceEntropyNats": {
            "mean": 0.38518914890588557,
            "median": 0.13944099177182165,
            "p95": 1.3614975180658118,
            "max": 1.9162432508737113
          },
          "top1Agreement": 0.7864583333333334
        }
      }
    },
    "officialReference": {
      "value": 0.022,
      "source": "Youssofal/Qwen3.8-27B-MTPLX-Optimized-Speed",
      "scope": "separate MTPLX coding battery, BF16 reference",
      "warning": "Not directly comparable to the local disclosed token-level suite."
    }
  }
}
