Files
Jyotisha/tests/test_rectification_relative_support.py
T
Jesse_ChenandClaude Opus 5.5 3733b9787b fix(rectification): bump scoring identity to scoring-10 for functional profile v2; dated contract by generation (BUG-1181)
Functional roles feed the *_functional_*_auxiliary rules, so 57782aea changes
candidate scores for identical input (memoization fixture 12:00: 8.6274 ->
8.4977). Per the "scoring semantics change => bump ALGORITHM_VERSION"
precedent (scoring-7 -> 8 -> 9), the identity moves to scoring-10; policy v3,
input contract v5 and Skill versions are unchanged, history is not relabeled.

Five frontend sites and one SQL guard tested `=== "...scoring-9"` for the
dated candidate-window contract; they now use isDatedScoringAlgorithmVersion /
a generation regex (>= 9). Migration 20261002010000 only recreates
validate_dated_rectification_candidate (one-line guard change).

Memoization golden v2 written by the test's own write_golden; v1 (scoring-8)
frozen by sha256. Real-engine scoring-10 cross-midnight golden added. Research
records re-frozen per ERR-110 (label functional_v2_2026_10_02) and
scripts/functional_benefics.py added to the frozen production identity
(ERR-114: 57782aea changed scores without tripping it).

Co-Authored-By: Claude Opus 5.5 (1M context) <noreply@anthropic.com>
Claude-Session: https://claude.ai/code/session_01N4f2nya58RoRu4yEmJgRGE
2026-10-02 12:31:50 +08:00

68 lines
3.2 KiB
Python

from __future__ import annotations
import unittest
from decimal import Decimal
from scripts.rectification.decision_policy import (
POLICY_VERSION,
RELATIVE_SUPPORT_MODE,
_relative_support,
build_candidate_decisions,
)
from scripts.rectification.scoring_service import ALGORITHM_VERSION
def _row(time: str, score: float) -> dict:
return {"time": time, "score": score, "evidence": [], "missing_layers": []}
class RelativeSupportScaleTest(unittest.TestCase):
def test_holdout_gate_keeps_proportional_default(self) -> None:
self.assertEqual(POLICY_VERSION, "rectification-candidate-policy-v3")
# Explicit date windows: scoring identity -8 -> -9, not a prior/policy change.
# 原值: -9;新值: -10;原因: 功能吉凶 v2 改变打分(BUG-1181),策略 v3 与相对支持口径不变。
self.assertEqual(ALGORITHM_VERSION, "rectification-v5-matrix-scoring-10")
self.assertEqual(RELATIVE_SUPPORT_MODE, "proportional")
def test_offset_top_two_lead_is_at_least_proportional(self) -> None:
scores = [
Decimal("15.52"), Decimal("15.40"), Decimal("14.10"), Decimal("13.00"),
Decimal("12.50"), Decimal("12.00"), Decimal("11.80"), Decimal("11.50"),
Decimal("11.40"), Decimal("11.20"), Decimal("11.10"), Decimal("10.96"),
]
floor = min(scores)
proportional = _relative_support(scores, floor=floor, mode="proportional")
offset = _relative_support(scores, floor=floor, mode="offset")
prop_lead = sorted(proportional, reverse=True)
offset_lead = sorted(offset, reverse=True)
self.assertGreaterEqual(offset_lead[0] - offset_lead[1], prop_lead[0] - prop_lead[1])
self.assertEqual(sum(offset), 100)
self.assertEqual(sum(proportional), 100)
def test_equal_scores_split_evenly_for_offset_and_proportional(self) -> None:
scores = [Decimal("12")] * 12
expected = _relative_support(scores, floor=Decimal("12"), mode="proportional")
offset = _relative_support(scores, floor=Decimal("12"), mode="offset")
self.assertEqual(offset, expected)
self.assertEqual(sum(offset), 100)
self.assertTrue(all(value in {8, 9} for value in offset))
def test_build_candidate_decisions_offset_lead_beats_proportional_on_spaced_grid(self) -> None:
scores = [
15.52, 15.40, 14.10, 13.00, 12.50, 12.00,
11.80, 11.50, 11.40, 11.20, 11.10, 10.96,
]
rows = [_row(f"05:{index * 5:02d}", score) for index, score in enumerate(scores)]
proportional = build_candidate_decisions(rows, result_id="00000000-0000-4000-8000-000000000001", support_mode="proportional")
offset = build_candidate_decisions(rows, result_id="00000000-0000-4000-8000-000000000001", support_mode="offset")
self.assertEqual(len(proportional), 12)
self.assertEqual(len(offset), 12)
prop_lead = proportional[0]["relative_support"] - proportional[1]["relative_support"]
offset_lead = offset[0]["relative_support"] - offset[1]["relative_support"]
self.assertGreaterEqual(offset_lead, prop_lead)
self.assertEqual(sum(item["relative_support"] for item in offset), 100)
if __name__ == "__main__":
unittest.main()