{"abstract":"Seven-judge dive scores sum five marks and come out far too high.","category":"Sports scoring and tiebreakers","checks":8,"contract":"Diving dive score. scores are judge marks as strings from 0 to 10 in half points; any other mark returns \"invalid mark <s>\". A 5-judge panel drops the single highest and lowest marks, a 7-judge panel the two highest and two lowest; any other panel size returns \"invalid panel\". The three remaining marks are summed and multiplied by the degree of difficulty dd (a decimal string); return the exact result with two decimals.","evaluation_group":"w2-sports-scoring-diving-judges-trim","failed_approach":"Dropping three from each end keeps only the median mark.","family":"w2-sports-scoring-diving-judges-trim-seven-judge-trim","id":"FA-84246","implementations":{"attempt":{"sha256":"2253001aa712b6d40883bbb4772dc59a710f3a44807613f37bb70788c127e74c","source":"\"\"\"Failure Map reference implementation. Python standard library only.\"\"\"\nimport json\nfrom fractions import Fraction\nN = 1\nobservations = []\ndef solve(scores, dd):\n    marks = []\n    for s in scores:\n        v = Fraction(s)\n        if v < 0 or v > 10 or (v * 2).denominator != 1:\n            return 'invalid mark ' + s\n        marks.append(v)\n    if len(marks) == 5:\n        drop = 1\n    elif len(marks) == 7:\n        drop = 3\n    else:\n        return 'invalid panel'\n    kept = sorted(marks)[drop:len(marks) - drop]\n    total = sum(kept) * Fraction(dd)\n    return '%.2f' % float(total)\ndef check(label, actual, expected):\n    observations.append({\"check\": label, \"actual\": actual, \"expected\": expected, \"passed\": actual == expected})\ndef run(args):\n    try:\n        return solve(*args)\n    except Exception as exc:\n        return 'raised ' + type(exc).__name__\ncases = [[('control five-judge panel', (['7.0', '7.5', '6.5', '8.0', '7.0'], '2.0'), '43.00'),\n  ('control seven-judge panel',\n   (['8.0', '8.5', '9.0', '7.5', '8.0', '8.5', '6.0'], '3.1'),\n   '75.95'),\n  ('boundary perfect tens', (['10.0', '10.0', '10.0', '10.0', '10.0'], '3.4'), '102.00'),\n  ('boundary non-half mark', (['7.3', '7.0', '7.0', '7.0', '7.0'], '2.0'), 'invalid mark 7.3'),\n  ('boundary six judges', (['7.0', '7.0', '7.0', '7.0', '7.0', '7.0'], '2.0'), 'invalid panel'),\n  ('regression: seven judge trim',\n   (['5.5', '1.5', '4.0', '1.5', '7.5', '0.5', '9.0'], '3.7'),\n   '40.70'),\n  ('variant scenario 1', (['8.5', '4.5', '2.5', '3.5', '5.5'], '3.4'), '45.90'),\n  ('variant scenario 2', (['4.0', '3.0', '8.0', '8.5', '1.0'], '2.0'), '30.00')],\n [('control five-judge panel', (['7.0', '7.5', '6.5', '8.0', '7.0'], '2.0'), '43.00'),\n  ('control seven-judge panel',\n   (['8.0', '8.5', '9.0', '7.5', '8.0', '8.5', '6.0'], '3.1'),\n   '75.95'),\n  ('boundary perfect tens', (['10.0', '10.0', '10.0', '10.0', '10.0'], '3.4'), '102.00'),\n  ('boundary non-half mark', (['7.3', '7.0', '7.0', '7.0', '7.0'], '2.0'), 'invalid mark 7.3'),\n  ('boundary six judges', (['7.0', '7.0', '7.0', '7.0', '7.0', '7.0'], '2.0'), 'invalid panel'),\n  ('regression: seven judge trim',\n   (['3.5', '9.0', '5.5', '2.5', '3.0', '5.5', '4.0'], '3.4'),\n   '44.20'),\n  ('variant scenario 1',\n   (['7.3', '4.5', '4.5', '6.0', '0.0', '5.0', '0.5'], '3.1'),\n   'invalid mark 7.3'),\n  ('variant scenario 2', (['10.0', '8.0', '3.0', '3.0', '9.0'], '3.7'), '74.00')],\n [('control five-judge panel', (['7.0', '7.5', '6.5', '8.0', '7.0'], '2.0'), '43.00'),\n  ('control seven-judge panel',\n   (['8.0', '8.5', '9.0', '7.5', '8.0', '8.5', '6.0'], '3.1'),\n   '75.95'),\n  ('boundary perfect tens', (['10.0', '10.0', '10.0', '10.0', '10.0'], '3.4'), '102.00'),\n  ('boundary non-half mark', (['7.3', '7.0', '7.0', '7.0', '7.0'], '2.0'), 'invalid mark 7.3'),\n  ('boundary six judges', (['7.0', '7.0', '7.0', '7.0', '7.0', '7.0'], '2.0'), 'invalid panel'),\n  ('regression: seven judge trim',\n   (['2.5', '6.5', '9.0', '3.5', '6.0', '0.5', '9.0'], '3.7'),\n   '59.20'),\n  ('variant scenario 1', (['1.5', '6.5', '10.0', '2.5', '8.0', '8.5'], '3.1'), 'invalid panel'),\n  ('variant scenario 2', (['8.0', '0.0', '6.5', '9.0', '5.0', '4.0', '6.0'], '3.7'), '64.75')],\n [('control five-judge panel', (['7.0', '7.5', '6.5', '8.0', '7.0'], '2.0'), '43.00'),\n  ('control seven-judge panel',\n   (['8.0', '8.5', '9.0', '7.5', '8.0', '8.5', '6.0'], '3.1'),\n   '75.95'),\n  ('boundary perfect tens', (['10.0', '10.0', '10.0', '10.0', '10.0'], '3.4'), '102.00'),\n  ('boundary non-half mark', (['7.3', '7.0', '7.0', '7.0', '7.0'], '2.0'), 'invalid mark 7.3'),\n  ('boundary six judges', (['7.0', '7.0', '7.0', '7.0', '7.0', '7.0'], '2.0'), 'invalid panel'),\n  ('regression: seven judge trim',\n   (['7.5', '7.0', '10.0', '0.0', '1.5', '2.5', '9.0'], '1.6'),\n   '27.20'),\n  ('variant scenario 1', (['10.0', '2.5', '7.0', '1.0', '9.5'], '1.6'), '30.40'),\n  ('variant scenario 2', (['1.5', '0.5', '7.5', '8.0', '7.0', '1.0'], '3.7'), 'invalid panel')],\n [('control five-judge panel', (['7.0', '7.5', '6.5', '8.0', '7.0'], '2.0'), '43.00'),\n  ('control seven-judge panel',\n   (['8.0', '8.5', '9.0', '7.5', '8.0', '8.5', '6.0'], '3.1'),\n   '75.95'),\n  ('boundary perfect tens', (['10.0', '10.0', '10.0', '10.0', '10.0'], '3.4'), '102.00'),\n  ('boundary non-half mark', (['7.3', '7.0', '7.0', '7.0', '7.0'], '2.0'), 'invalid mark 7.3'),\n  ('boundary six judges', (['7.0', '7.0', '7.0', '7.0', '7.0', '7.0'], '2.0'), 'invalid panel'),\n  ('regression: seven judge trim',\n   (['1.5', '10.0', '6.0', '4.0', '9.0', '4.5', '0.0'], '3.1'),\n   '44.95'),\n  ('variant scenario 1', (['4.5', '5.5', '1.5', '10.0', '9.5', '9.5'], '2.8'), 'invalid panel'),\n  ('variant scenario 2', (['8.5', '5.5', '4.0', '0.5', '10.0'], '3.4'), '61.20')]]\nfor label, args, expected in cases[N - 1]:\n    check(label, run(args), expected)\nprint(json.dumps({\"observations\": observations, \"passed\": all(x[\"passed\"] for x in observations)}, ensure_ascii=False))\nraise SystemExit(0 if all(x[\"passed\"] for x in observations) else 1)\n"},"broken":{"sha256":"68ab8c63011d35ec09c52729e6ac8370195978c03b346fc5098f05d3de0fc7af","source":"\"\"\"Failure Map reference implementation. Python standard library only.\"\"\"\nimport json\nfrom fractions import Fraction\nN = 1\nobservations = []\ndef solve(scores, dd):\n    marks = []\n    for s in scores:\n        v = Fraction(s)\n        if v < 0 or v > 10 or (v * 2).denominator != 1:\n            return 'invalid mark ' + s\n        marks.append(v)\n    if len(marks) == 5:\n        drop = 1\n    elif len(marks) == 7:\n        drop = 1\n    else:\n        return 'invalid panel'\n    kept = sorted(marks)[drop:len(marks) - drop]\n    total = sum(kept) * Fraction(dd)\n    return '%.2f' % float(total)\ndef check(label, actual, expected):\n    observations.append({\"check\": label, \"actual\": actual, \"expected\": expected, \"passed\": actual == expected})\ndef run(args):\n    try:\n        return solve(*args)\n    except Exception as exc:\n        return 'raised ' + type(exc).__name__\ncases = [[('control five-judge panel', (['7.0', '7.5', '6.5', '8.0', '7.0'], '2.0'), '43.00'),\n  ('control seven-judge panel',\n   (['8.0', '8.5', '9.0', '7.5', '8.0', '8.5', '6.0'], '3.1'),\n   '75.95'),\n  ('boundary perfect tens', (['10.0', '10.0', '10.0', '10.0', '10.0'], '3.4'), '102.00'),\n  ('boundary non-half mark', (['7.3', '7.0', '7.0', '7.0', '7.0'], '2.0'), 'invalid mark 7.3'),\n  ('boundary six judges', (['7.0', '7.0', '7.0', '7.0', '7.0', '7.0'], '2.0'), 'invalid panel'),\n  ('regression: seven judge trim',\n   (['5.5', '1.5', '4.0', '1.5', '7.5', '0.5', '9.0'], '3.7'),\n   '40.70'),\n  ('variant scenario 1', (['8.5', '4.5', '2.5', '3.5', '5.5'], '3.4'), '45.90'),\n  ('variant scenario 2', (['4.0', '3.0', '8.0', '8.5', '1.0'], '2.0'), '30.00')],\n [('control five-judge panel', (['7.0', '7.5', '6.5', '8.0', '7.0'], '2.0'), '43.00'),\n  ('control seven-judge panel',\n   (['8.0', '8.5', '9.0', '7.5', '8.0', '8.5', '6.0'], '3.1'),\n   '75.95'),\n  ('boundary perfect tens', (['10.0', '10.0', '10.0', '10.0', '10.0'], '3.4'), '102.00'),\n  ('boundary non-half mark', (['7.3', '7.0', '7.0', '7.0', '7.0'], '2.0'), 'invalid mark 7.3'),\n  ('boundary six judges', (['7.0', '7.0', '7.0', '7.0', '7.0', '7.0'], '2.0'), 'invalid panel'),\n  ('regression: seven judge trim',\n   (['3.5', '9.0', '5.5', '2.5', '3.0', '5.5', '4.0'], '3.4'),\n   '44.20'),\n  ('variant scenario 1',\n   (['7.3', '4.5', '4.5', '6.0', '0.0', '5.0', '0.5'], '3.1'),\n   'invalid mark 7.3'),\n  ('variant scenario 2', (['10.0', '8.0', '3.0', '3.0', '9.0'], '3.7'), '74.00')],\n [('control five-judge panel', (['7.0', '7.5', '6.5', '8.0', '7.0'], '2.0'), '43.00'),\n  ('control seven-judge panel',\n   (['8.0', '8.5', '9.0', '7.5', '8.0', '8.5', '6.0'], '3.1'),\n   '75.95'),\n  ('boundary perfect tens', (['10.0', '10.0', '10.0', '10.0', '10.0'], '3.4'), '102.00'),\n  ('boundary non-half mark', (['7.3', '7.0', '7.0', '7.0', '7.0'], '2.0'), 'invalid mark 7.3'),\n  ('boundary six judges', (['7.0', '7.0', '7.0', '7.0', '7.0', '7.0'], '2.0'), 'invalid panel'),\n  ('regression: seven judge trim',\n   (['2.5', '6.5', '9.0', '3.5', '6.0', '0.5', '9.0'], '3.7'),\n   '59.20'),\n  ('variant scenario 1', (['1.5', '6.5', '10.0', '2.5', '8.0', '8.5'], '3.1'), 'invalid panel'),\n  ('variant scenario 2', (['8.0', '0.0', '6.5', '9.0', '5.0', '4.0', '6.0'], '3.7'), '64.75')],\n [('control five-judge panel', (['7.0', '7.5', '6.5', '8.0', '7.0'], '2.0'), '43.00'),\n  ('control seven-judge panel',\n   (['8.0', '8.5', '9.0', '7.5', '8.0', '8.5', '6.0'], '3.1'),\n   '75.95'),\n  ('boundary perfect tens', (['10.0', '10.0', '10.0', '10.0', '10.0'], '3.4'), '102.00'),\n  ('boundary non-half mark', (['7.3', '7.0', '7.0', '7.0', '7.0'], '2.0'), 'invalid mark 7.3'),\n  ('boundary six judges', (['7.0', '7.0', '7.0', '7.0', '7.0', '7.0'], '2.0'), 'invalid panel'),\n  ('regression: seven judge trim',\n   (['7.5', '7.0', '10.0', '0.0', '1.5', '2.5', '9.0'], '1.6'),\n   '27.20'),\n  ('variant scenario 1', (['10.0', '2.5', '7.0', '1.0', '9.5'], '1.6'), '30.40'),\n  ('variant scenario 2', (['1.5', '0.5', '7.5', '8.0', '7.0', '1.0'], '3.7'), 'invalid panel')],\n [('control five-judge panel', (['7.0', '7.5', '6.5', '8.0', '7.0'], '2.0'), '43.00'),\n  ('control seven-judge panel',\n   (['8.0', '8.5', '9.0', '7.5', '8.0', '8.5', '6.0'], '3.1'),\n   '75.95'),\n  ('boundary perfect tens', (['10.0', '10.0', '10.0', '10.0', '10.0'], '3.4'), '102.00'),\n  ('boundary non-half mark', (['7.3', '7.0', '7.0', '7.0', '7.0'], '2.0'), 'invalid mark 7.3'),\n  ('boundary six judges', (['7.0', '7.0', '7.0', '7.0', '7.0', '7.0'], '2.0'), 'invalid panel'),\n  ('regression: seven judge trim',\n   (['1.5', '10.0', '6.0', '4.0', '9.0', '4.5', '0.0'], '3.1'),\n   '44.95'),\n  ('variant scenario 1', (['4.5', '5.5', '1.5', '10.0', '9.5', '9.5'], '2.8'), 'invalid panel'),\n  ('variant scenario 2', (['8.5', '5.5', '4.0', '0.5', '10.0'], '3.4'), '61.20')]]\nfor label, args, expected in cases[N - 1]:\n    check(label, run(args), expected)\nprint(json.dumps({\"observations\": observations, \"passed\": all(x[\"passed\"] for x in observations)}, ensure_ascii=False))\nraise SystemExit(0 if all(x[\"passed\"] for x in observations) else 1)\n"},"fixed":{"sha256":"6e3b06cd020371d144704e782842d94c8467c2bdaddcaaeb20a51078d107f2e9","source":"\"\"\"Failure Map reference implementation. Python standard library only.\"\"\"\nimport json\nfrom fractions import Fraction\nN = 1\nobservations = []\ndef solve(scores, dd):\n    marks = []\n    for s in scores:\n        v = Fraction(s)\n        if v < 0 or v > 10 or (v * 2).denominator != 1:\n            return 'invalid mark ' + s\n        marks.append(v)\n    if len(marks) == 5:\n        drop = 1\n    elif len(marks) == 7:\n        drop = 2\n    else:\n        return 'invalid panel'\n    kept = sorted(marks)[drop:len(marks) - drop]\n    total = sum(kept) * Fraction(dd)\n    return '%.2f' % float(total)\ndef check(label, actual, expected):\n    observations.append({\"check\": label, \"actual\": actual, \"expected\": expected, \"passed\": actual == expected})\ndef run(args):\n    try:\n        return solve(*args)\n    except Exception as exc:\n        return 'raised ' + type(exc).__name__\ncases = [[('control five-judge panel', (['7.0', '7.5', '6.5', '8.0', '7.0'], '2.0'), '43.00'),\n  ('control seven-judge panel',\n   (['8.0', '8.5', '9.0', '7.5', '8.0', '8.5', '6.0'], '3.1'),\n   '75.95'),\n  ('boundary perfect tens', (['10.0', '10.0', '10.0', '10.0', '10.0'], '3.4'), '102.00'),\n  ('boundary non-half mark', (['7.3', '7.0', '7.0', '7.0', '7.0'], '2.0'), 'invalid mark 7.3'),\n  ('boundary six judges', (['7.0', '7.0', '7.0', '7.0', '7.0', '7.0'], '2.0'), 'invalid panel'),\n  ('regression: seven judge trim',\n   (['5.5', '1.5', '4.0', '1.5', '7.5', '0.5', '9.0'], '3.7'),\n   '40.70'),\n  ('variant scenario 1', (['8.5', '4.5', '2.5', '3.5', '5.5'], '3.4'), '45.90'),\n  ('variant scenario 2', (['4.0', '3.0', '8.0', '8.5', '1.0'], '2.0'), '30.00')],\n [('control five-judge panel', (['7.0', '7.5', '6.5', '8.0', '7.0'], '2.0'), '43.00'),\n  ('control seven-judge panel',\n   (['8.0', '8.5', '9.0', '7.5', '8.0', '8.5', '6.0'], '3.1'),\n   '75.95'),\n  ('boundary perfect tens', (['10.0', '10.0', '10.0', '10.0', '10.0'], '3.4'), '102.00'),\n  ('boundary non-half mark', (['7.3', '7.0', '7.0', '7.0', '7.0'], '2.0'), 'invalid mark 7.3'),\n  ('boundary six judges', (['7.0', '7.0', '7.0', '7.0', '7.0', '7.0'], '2.0'), 'invalid panel'),\n  ('regression: seven judge trim',\n   (['3.5', '9.0', '5.5', '2.5', '3.0', '5.5', '4.0'], '3.4'),\n   '44.20'),\n  ('variant scenario 1',\n   (['7.3', '4.5', '4.5', '6.0', '0.0', '5.0', '0.5'], '3.1'),\n   'invalid mark 7.3'),\n  ('variant scenario 2', (['10.0', '8.0', '3.0', '3.0', '9.0'], '3.7'), '74.00')],\n [('control five-judge panel', (['7.0', '7.5', '6.5', '8.0', '7.0'], '2.0'), '43.00'),\n  ('control seven-judge panel',\n   (['8.0', '8.5', '9.0', '7.5', '8.0', '8.5', '6.0'], '3.1'),\n   '75.95'),\n  ('boundary perfect tens', (['10.0', '10.0', '10.0', '10.0', '10.0'], '3.4'), '102.00'),\n  ('boundary non-half mark', (['7.3', '7.0', '7.0', '7.0', '7.0'], '2.0'), 'invalid mark 7.3'),\n  ('boundary six judges', (['7.0', '7.0', '7.0', '7.0', '7.0', '7.0'], '2.0'), 'invalid panel'),\n  ('regression: seven judge trim',\n   (['2.5', '6.5', '9.0', '3.5', '6.0', '0.5', '9.0'], '3.7'),\n   '59.20'),\n  ('variant scenario 1', (['1.5', '6.5', '10.0', '2.5', '8.0', '8.5'], '3.1'), 'invalid panel'),\n  ('variant scenario 2', (['8.0', '0.0', '6.5', '9.0', '5.0', '4.0', '6.0'], '3.7'), '64.75')],\n [('control five-judge panel', (['7.0', '7.5', '6.5', '8.0', '7.0'], '2.0'), '43.00'),\n  ('control seven-judge panel',\n   (['8.0', '8.5', '9.0', '7.5', '8.0', '8.5', '6.0'], '3.1'),\n   '75.95'),\n  ('boundary perfect tens', (['10.0', '10.0', '10.0', '10.0', '10.0'], '3.4'), '102.00'),\n  ('boundary non-half mark', (['7.3', '7.0', '7.0', '7.0', '7.0'], '2.0'), 'invalid mark 7.3'),\n  ('boundary six judges', (['7.0', '7.0', '7.0', '7.0', '7.0', '7.0'], '2.0'), 'invalid panel'),\n  ('regression: seven judge trim',\n   (['7.5', '7.0', '10.0', '0.0', '1.5', '2.5', '9.0'], '1.6'),\n   '27.20'),\n  ('variant scenario 1', (['10.0', '2.5', '7.0', '1.0', '9.5'], '1.6'), '30.40'),\n  ('variant scenario 2', (['1.5', '0.5', '7.5', '8.0', '7.0', '1.0'], '3.7'), 'invalid panel')],\n [('control five-judge panel', (['7.0', '7.5', '6.5', '8.0', '7.0'], '2.0'), '43.00'),\n  ('control seven-judge panel',\n   (['8.0', '8.5', '9.0', '7.5', '8.0', '8.5', '6.0'], '3.1'),\n   '75.95'),\n  ('boundary perfect tens', (['10.0', '10.0', '10.0', '10.0', '10.0'], '3.4'), '102.00'),\n  ('boundary non-half mark', (['7.3', '7.0', '7.0', '7.0', '7.0'], '2.0'), 'invalid mark 7.3'),\n  ('boundary six judges', (['7.0', '7.0', '7.0', '7.0', '7.0', '7.0'], '2.0'), 'invalid panel'),\n  ('regression: seven judge trim',\n   (['1.5', '10.0', '6.0', '4.0', '9.0', '4.5', '0.0'], '3.1'),\n   '44.95'),\n  ('variant scenario 1', (['4.5', '5.5', '1.5', '10.0', '9.5', '9.5'], '2.8'), 'invalid panel'),\n  ('variant scenario 2', (['8.5', '5.5', '4.0', '0.5', '10.0'], '3.4'), '61.20')]]\nfor label, args, expected in cases[N - 1]:\n    check(label, run(args), expected)\nprint(json.dumps({\"observations\": observations, \"passed\": all(x[\"passed\"] for x in observations)}, ensure_ascii=False))\nraise SystemExit(0 if all(x[\"passed\"] for x in observations) else 1)\n"}},"limitations":"Stipulated, bounded toy contract stated in the contract field; not a claim of conformance with any governing body rulebook or operator house rules. This reproducer isolates one failure mechanism. Results cover the supplied fixtures. Variants within a family share a test contract and should remain grouped when constructing evaluation splits. Related mechanisms with a shared evaluation_group must also remain together; these controlled models are not independent production incidents.","method":"Deterministic executable model with adversarial boundary fixtures.","provenance":{"created_by":"Failure Map","dependencies":"Python standard library","family":"w2-sports-scoring-diving-judges-trim-seven-judge-trim","generated_at":"2026-09-29T14:50:29.144522+00:00","license":"CC0-1.0","python":"3.12.14","seed":1,"split":"open-access"},"relevance":"Meet management systems compute dive scores from judge panels of different sizes.","repair":"Drop two high and two low marks for seven judges.","root_cause":"The drop count for seven judges is 1.","sha256":"b238622b9126f4cae2b482eb1c9b9a7bcac86700ce45887c50cc8be57a4f6040","title":"Seven-judge panel trimmed like a five-judge panel · case 01","variant":1,"variant_policy":"Five numbered records share a model and may reuse boundary fixtures.","verification":{"attempt":{"elapsed_ms":42.538,"exit_code":1,"observations":[{"actual":"43.00","check":"control five-judge panel","expected":"43.00","passed":true},{"actual":"24.80","check":"control seven-judge panel","expected":"75.95","passed":false},{"actual":"102.00","check":"boundary perfect tens","expected":"102.00","passed":true},{"actual":"invalid mark 7.3","check":"boundary non-half mark","expected":"invalid mark 7.3","passed":true},{"actual":"invalid panel","check":"boundary six judges","expected":"invalid panel","passed":true},{"actual":"14.80","check":"regression: seven judge trim","expected":"40.70","passed":false},{"actual":"45.90","check":"variant scenario 1","expected":"45.90","passed":true},{"actual":"30.00","check":"variant scenario 2","expected":"30.00","passed":true}],"passed":false,"stderr":"","stdout":"{\"observations\": [{\"check\": \"control five-judge panel\", \"actual\": \"43.00\", \"expected\": \"43.00\", \"passed\": true}, {\"check\": \"control seven-judge panel\", \"actual\": \"24.80\", \"expected\": \"75.95\", \"passed\": false}, {\"check\": \"boundary perfect tens\", \"actual\": \"102.00\", \"expected\": \"102.00\", \"passed\": true}, {\"check\": \"boundary non-half mark\", \"actual\": \"invalid mark 7.3\", \"expected\": \"invalid mark 7.3\", \"passed\": true}, {\"check\": \"boundary six judges\", \"actual\": \"invalid panel\", \"expected\": \"invalid panel\", \"passed\": true}, {\"check\": \"regression: seven judge trim\", \"actual\": \"14.80\", \"expected\": \"40.70\", \"passed\": false}, {\"check\": \"variant scenario 1\", \"actual\": \"45.90\", \"expected\": \"45.90\", \"passed\": true}, {\"check\": \"variant scenario 2\", \"actual\": \"30.00\", \"expected\": \"30.00\", \"passed\": true}], \"passed\": false}\n"},"broken":{"elapsed_ms":43.845,"exit_code":1,"observations":[{"actual":"43.00","check":"control five-judge panel","expected":"43.00","passed":true},{"actual":"125.55","check":"control seven-judge panel","expected":"75.95","passed":false},{"actual":"102.00","check":"boundary perfect tens","expected":"102.00","passed":true},{"actual":"invalid mark 7.3","check":"boundary non-half mark","expected":"invalid mark 7.3","passed":true},{"actual":"invalid panel","check":"boundary six judges","expected":"invalid panel","passed":true},{"actual":"74.00","check":"regression: seven judge trim","expected":"40.70","passed":false},{"actual":"45.90","check":"variant scenario 1","expected":"45.90","passed":true},{"actual":"30.00","check":"variant scenario 2","expected":"30.00","passed":true}],"passed":false,"stderr":"","stdout":"{\"observations\": [{\"check\": \"control five-judge panel\", \"actual\": \"43.00\", \"expected\": \"43.00\", \"passed\": true}, {\"check\": \"control seven-judge panel\", \"actual\": \"125.55\", \"expected\": \"75.95\", \"passed\": false}, {\"check\": \"boundary perfect tens\", \"actual\": \"102.00\", \"expected\": \"102.00\", \"passed\": true}, {\"check\": \"boundary non-half mark\", \"actual\": \"invalid mark 7.3\", \"expected\": \"invalid mark 7.3\", \"passed\": true}, {\"check\": \"boundary six judges\", \"actual\": \"invalid panel\", \"expected\": \"invalid panel\", \"passed\": true}, {\"check\": \"regression: seven judge trim\", \"actual\": \"74.00\", \"expected\": \"40.70\", \"passed\": false}, {\"check\": \"variant scenario 1\", \"actual\": \"45.90\", \"expected\": \"45.90\", \"passed\": true}, {\"check\": \"variant scenario 2\", \"actual\": \"30.00\", \"expected\": \"30.00\", \"passed\": true}], \"passed\": false}\n"},"fixed":{"elapsed_ms":42.183,"exit_code":0,"observations":[{"actual":"43.00","check":"control five-judge panel","expected":"43.00","passed":true},{"actual":"75.95","check":"control seven-judge panel","expected":"75.95","passed":true},{"actual":"102.00","check":"boundary perfect tens","expected":"102.00","passed":true},{"actual":"invalid mark 7.3","check":"boundary non-half mark","expected":"invalid mark 7.3","passed":true},{"actual":"invalid panel","check":"boundary six judges","expected":"invalid panel","passed":true},{"actual":"40.70","check":"regression: seven judge trim","expected":"40.70","passed":true},{"actual":"45.90","check":"variant scenario 1","expected":"45.90","passed":true},{"actual":"30.00","check":"variant scenario 2","expected":"30.00","passed":true}],"passed":true,"stderr":"","stdout":"{\"observations\": [{\"check\": \"control five-judge panel\", \"actual\": \"43.00\", \"expected\": \"43.00\", \"passed\": true}, {\"check\": \"control seven-judge panel\", \"actual\": \"75.95\", \"expected\": \"75.95\", \"passed\": true}, {\"check\": \"boundary perfect tens\", \"actual\": \"102.00\", \"expected\": \"102.00\", \"passed\": true}, {\"check\": \"boundary non-half mark\", \"actual\": \"invalid mark 7.3\", \"expected\": \"invalid mark 7.3\", \"passed\": true}, {\"check\": \"boundary six judges\", \"actual\": \"invalid panel\", \"expected\": \"invalid panel\", \"passed\": true}, {\"check\": \"regression: seven judge trim\", \"actual\": \"40.70\", \"expected\": \"40.70\", \"passed\": true}, {\"check\": \"variant scenario 1\", \"actual\": \"45.90\", \"expected\": \"45.90\", \"passed\": true}, {\"check\": \"variant scenario 2\", \"actual\": \"30.00\", \"expected\": \"30.00\", \"passed\": true}], \"passed\": true}\n"}},"verified":true,"visibility":"public"}