diff --git a/bench/baseline.json b/bench/baseline.json index 14abc90..d9d5cfd 100644 --- a/bench/baseline.json +++ b/bench/baseline.json @@ -2,7 +2,7 @@ "version": 1, "captured": "2026-05-20", "captured_via": "bench/run.sh -n 5", - "note": "Baseline for bench/check.py regression detection. The language-invariant thresholds (rc/bump <= 1.3x throughput, p99/median <= 5x latency) are NOT the regression-check tolerances; the per-metric tolerances below are tuned to absorb run-to-run noise on a quiet developer machine. To update after an intentional change, re-run bench/run.sh and replace the values, recording the reason in the commit body that ships the baseline bump. The latency arms gate only on median / p99 / p99_over_median \u2014 max_us and p99_9_us were removed on the 2026-05-20 recapture (Gitea #15 / #16) because tail-of-distribution latency metrics are dominated by OS-level jitter (THP defrag, scheduler preemption, IRQ load), not allocator behaviour, and produced 3+ consecutive false-positive REGRESSION rows on byte-identical no-op milestones. See docs/specs/0047-bench-harness-recalibration.md.", + "note": "Baseline for bench/check.py regression detection. The language-invariant thresholds (rc/bump <= 1.3x throughput, p99/median <= 5x latency) are NOT the regression-check tolerances; the per-metric tolerances below are tuned to absorb run-to-run noise on a quiet developer machine. To update after an intentional change, re-run bench/run.sh and replace the values, recording the reason in the commit body that ships the baseline bump. The latency gate history: max_us and p99_9_us were removed on the 2026-05-20 recapture (Gitea #15 / #16) because tail-of-distribution latency metrics are dominated by OS-level jitter (THP defrag, scheduler preemption, IRQ load), not allocator behaviour, and produced 3+ consecutive false-positive REGRESSION rows on byte-identical no-op milestones; see docs/specs/0047-bench-harness-recalibration.md. At the 2026-05-28 kernel-extension-mechanics audit close, latency.explicit_at_rc.p99_us and latency.explicit_at_rc.p99_over_median were ALSO removed after a 6-invocation bencher characterisation found the explicit-arm p99 cv at 18.6% (vs implicit-arm cv 2.9%) with within-invocation 4-run spread up to 1.87x \u2014 structurally identical to the 2026-05-20 metric-removal precedent. The explicit-arm now gates only on median_us (which has cv 0.37% across 6 invocations, the actual allocator signal). The implicit-arm p99 / p99_over_median is kept because it is stable (cv 2.9%) and provides the contrast that lets future audits distinguish allocator regressions from machine jitter.", "throughput": { "bench_list_sum": { "bump_s": { @@ -142,14 +142,6 @@ "median_us": { "baseline": 218.8, "tolerance_pct": 15 - }, - "p99_us": { - "baseline": 259.9, - "tolerance_pct": 25 - }, - "p99_over_median": { - "baseline": 1.19, - "tolerance_pct": 25 } }, "implicit_at_rc": { diff --git a/examples/bench_latency_explicit.ail b/examples/bench_latency_explicit.ail index 6c95728..10e71b2 100644 --- a/examples/bench_latency_explicit.ail +++ b/examples/bench_latency_explicit.ail @@ -139,10 +139,10 @@ (params remaining print_countdown chunk_len print_k t) (body (if (app eq remaining 0) - (app print 9999) + (seq (app print 9999) (do io/print_str "\n")) (if (app eq print_countdown 0) (seq - (app print (app one_op chunk_len t)) + (seq (app print (app one_op chunk_len t)) (do io/print_str "\n")) (tail-app loop (app - remaining 1) (app - print_k 1) @@ -165,5 +165,5 @@ (let t (app build_tree 19) (let _root (app pin_root t) (seq - (app print 8888) + (seq (app print 8888) (do io/print_str "\n")) (app loop 20000 0 500 20 t))))))) diff --git a/examples/bench_latency_implicit.ail b/examples/bench_latency_implicit.ail index d25818b..c62edd9 100644 --- a/examples/bench_latency_implicit.ail +++ b/examples/bench_latency_implicit.ail @@ -198,10 +198,10 @@ (params remaining print_countdown chunk_len print_k t) (body (if (app eq remaining 0) - (app print 9999) + (seq (app print 9999) (do io/print_str "\n")) (if (app eq print_countdown 0) (seq - (app print (app one_op chunk_len t)) + (seq (app print (app one_op chunk_len t)) (do io/print_str "\n")) (tail-app loop (app - remaining 1) (app - print_k 1) @@ -224,5 +224,5 @@ (let t (app build_tree 19) (let _root (app pin_root t) (seq - (app print 8888) + (seq (app print 8888) (do io/print_str "\n")) (app loop 20000 0 500 20 t)))))))