forked from ianstormtaylor/slate
-
Notifications
You must be signed in to change notification settings - Fork 1
Expand file tree
/
Copy pathautoresearch.jsonl
More file actions
21 lines (21 loc) · 43.2 KB
/
Copy pathautoresearch.jsonl
File metadata and controls
21 lines (21 loc) · 43.2 KB
1
2
3
4
5
6
7
8
9
10
11
12
13
14
15
16
17
18
19
20
21
{"type":"config","name":"react-huge-document-full","goal":"","metricName":"react_huge_doc_full_max_budget_ratio","metricUnit":"ratio","bestDirection":"lower"}
{"run":1,"commit":"","metric":1.82,"metrics":{"react_huge_doc_legacy_compare_worst_p95_ratio":0.35,"react_huge_doc_legacy_compare_worst_p95_delta_ms":-52.74,"react_huge_doc_type_to_paint_p95_ms":30.4,"react_huge_doc_dom_nodes_p95":303,"react_huge_doc_heap_mb_p95":16.31,"react_huge_doc_long_task_max_p95_ms":0,"react_huge_doc_full_failure_count":0,"react_huge_doc_full_legacy_compare_worst_p95_ratio":0.35,"react_huge_doc_full_core_worst_p95_ratio":10.5,"react_huge_doc_full_type_to_paint_p95_ms":74.9,"react_huge_doc_full_virtualized_type_to_paint_p95_ms":30.4,"react_huge_doc_full_dom_nodes_p95":20511,"react_huge_doc_full_long_task_max_p95_ms":0},"metricEligible":false,"status":"measure","evidenceStatus":"accepted","description":"Fresh huge-document full baseline: max budget ratio 1.82 with checks green","timestamp":1780337531000,"segment":0,"confidence":null,"packetFingerprint":"d58f138822265cdfa903a7e8c2c6a4b70ddb24e7b93a89dd56befc0f598b02d9","promotion":{"label":"measurement","reasons":["Logged as measure; metric evidence is trend-only and not finalizer evidence."]},"asi":{"hypothesis":"A clean react-huge-document-full session should optimize the aggregate huge-document suite, not the old history or pagination ledger.","evidence":"packet 1: react_huge_doc_full_max_budget_ratio=1.82, failure_count=0, core_worst_p95_ratio=10.5, type_to_paint_p95_ms=74.9, virtualized_type_to_paint_p95_ms=30.4, bun check passed","lane":"react-huge-document-full","family":"fresh-baseline","next_action_hint":"Scout the core_worst_p95_ratio=10.5 lane first; browser type-to-paint is still above the stretch budget but the aggregate max is now dominated by core huge-doc operations."},"benchmarkContract":{"command":"bash ./autoresearch.sh","checksCommand":"bash ./autoresearch.checks.sh","commandFile":"","envFile":"","surfaceHash":"168565aa676bf18c5ca56f781d6ff1ec47fe9bce2964216dbbfd959fa59389c2","files":[{"path":"autoresearch.sh","hash":"6d3d1a67661ec440d965ed38f0db82252434129a1b335e6ef961febd7c8e2bcd"},{"path":"autoresearch.checks.sh","hash":"9261aeaf280f5f66388ebfe2bea60e5dfb8d37ba8c69af4e608e7140f6fcdc20"},{"path":"package.json","hash":"4820e8774a8ee0fb005430a8af30b37f351200eb439460b10cd2ac47676252c2"}],"capturedAt":"2026-06-01T18:11:53.324Z"}}
{"run":2,"commit":"","metric":1.01,"metrics":{"react_huge_doc_legacy_compare_worst_p95_ratio":0.45,"react_huge_doc_legacy_compare_worst_p95_delta_ms":-21.91,"react_huge_doc_type_to_paint_p95_ms":32.5,"react_huge_doc_dom_nodes_p95":303,"react_huge_doc_heap_mb_p95":16.31,"react_huge_doc_long_task_max_p95_ms":0,"react_huge_doc_full_failure_count":0,"react_huge_doc_full_legacy_compare_worst_p95_ratio":0.45,"react_huge_doc_full_core_worst_p95_ratio":18,"react_huge_doc_full_type_to_paint_p95_ms":75.8,"react_huge_doc_full_virtualized_type_to_paint_p95_ms":32.5,"react_huge_doc_full_dom_nodes_p95":20511,"react_huge_doc_full_long_task_max_p95_ms":0},"metricEligible":false,"status":"measure","evidenceStatus":"accepted","description":"Partial-DOM promotion and selection export optimization reduced huge-document max budget ratio to 1.01 with checks green","timestamp":1780339968603,"segment":0,"confidence":null,"packetFingerprint":"8be2bfab622a0a6c9622670809266267b1bf071c6a0839695883a3fa9144f979","promotion":{"label":"measurement","reasons":["Logged as measure; metric evidence is trend-only and not finalizer evidence."]},"asi":{"hypothesis":"Focus placeholders before promotion, defer promoted selection until mounted DOM, runtime-scope placeholder previews, and share fast DOM range resolution to cut partial-DOM promotion latency without selection regressions","evidence":"react_huge_doc_full_max_budget_ratio=1.01; partialDOMPromotionP95Ms=43.68/50; browserTraceTypeToPaintP95Ms=75.8/75; bun check passed","lane":"react-huge-document-full","family":"partial-dom-promotion","next_action_hint":"Rerun or optimize browserTraceTypeToPaintP95Ms, now the only budget row above 1.0"},"benchmarkContract":{"command":"bash ./autoresearch.sh","checksCommand":"bash ./autoresearch.checks.sh","commandFile":"","envFile":"","surfaceHash":"168565aa676bf18c5ca56f781d6ff1ec47fe9bce2964216dbbfd959fa59389c2","files":[{"path":"autoresearch.sh","hash":"6d3d1a67661ec440d965ed38f0db82252434129a1b335e6ef961febd7c8e2bcd"},{"path":"autoresearch.checks.sh","hash":"9261aeaf280f5f66388ebfe2bea60e5dfb8d37ba8c69af4e608e7140f6fcdc20"},{"path":"package.json","hash":"4820e8774a8ee0fb005430a8af30b37f351200eb439460b10cd2ac47676252c2"}],"capturedAt":"2026-06-01T18:52:11.239Z"}}
{"run":3,"commit":"","metric":1.2,"metrics":{"react_huge_doc_legacy_compare_worst_p95_ratio":0.77,"react_huge_doc_legacy_compare_worst_p95_delta_ms":-8.77,"react_huge_doc_type_to_paint_p95_ms":32,"react_huge_doc_dom_nodes_p95":303,"react_huge_doc_heap_mb_p95":16.31,"react_huge_doc_long_task_max_p95_ms":0,"react_huge_doc_full_failure_count":0,"react_huge_doc_full_legacy_compare_worst_p95_ratio":0.77,"react_huge_doc_full_core_worst_p95_ratio":9.5,"react_huge_doc_full_type_to_paint_p95_ms":89.9,"react_huge_doc_full_virtualized_type_to_paint_p95_ms":32,"react_huge_doc_full_dom_nodes_p95":20510,"react_huge_doc_full_long_task_max_p95_ms":0},"metricEligible":false,"status":"measure","evidenceStatus":"accepted","description":"Repeat huge-document full packet shows remaining budget miss is staged browser type-to-paint variance, not partial-DOM promotion","timestamp":1780340138386,"segment":0,"confidence":null,"packetFingerprint":"96635cda229fb033eb4844093f96ed6afcd2ded6910ab194a2192f409af4ca44","promotion":{"label":"measurement","reasons":["Logged as measure; metric evidence is trend-only and not finalizer evidence."]},"asi":{"hypothesis":"The partial-DOM promotion fix holds, but staged browser type-to-paint is now the dominant unstable row","evidence":"react_huge_doc_full_max_budget_ratio=1.20; browserTraceTypeToPaintP95Ms=89.9/75; virtualizedTypeToPaintP95Ms=32/75; checks passed","lane":"react-huge-document-full","family":"browser-type-to-paint","next_action_hint":"Inspect staged browser trace artifact and optimize defaultAuto/stagedDomPresent type-to-paint"},"benchmarkContract":{"command":"bash ./autoresearch.sh","checksCommand":"bash ./autoresearch.checks.sh","commandFile":"","envFile":"","surfaceHash":"168565aa676bf18c5ca56f781d6ff1ec47fe9bce2964216dbbfd959fa59389c2","files":[{"path":"autoresearch.sh","hash":"6d3d1a67661ec440d965ed38f0db82252434129a1b335e6ef961febd7c8e2bcd"},{"path":"autoresearch.checks.sh","hash":"9261aeaf280f5f66388ebfe2bea60e5dfb8d37ba8c69af4e608e7140f6fcdc20"},{"path":"package.json","hash":"4820e8774a8ee0fb005430a8af30b37f351200eb439460b10cd2ac47676252c2"}],"capturedAt":"2026-06-01T18:55:28.047Z"}}
{"run":4,"commit":"7eaafb9","metric":1.06,"metrics":{"react_huge_doc_type_to_paint_p95_ms":26.4,"react_huge_doc_dom_nodes_p95":20510,"react_huge_doc_heap_mb_p95":73.05,"react_huge_doc_long_task_max_p95_ms":0,"react_huge_doc_full_failure_count":2,"react_huge_doc_full_legacy_compare_worst_p95_ratio":0.77,"react_huge_doc_full_core_worst_p95_ratio":11,"react_huge_doc_full_type_to_paint_p95_ms":26.4,"react_huge_doc_full_burst_to_paint_p95_ms":139.6,"react_huge_doc_full_virtualized_type_to_paint_p95_ms":26.4,"react_huge_doc_full_dom_nodes_p95":20510,"react_huge_doc_full_long_task_max_p95_ms":0},"metricEligible":false,"status":"crash","evidenceStatus":"rejected","description":"Full huge-document packet parsed metrics but exited 1 from browser trace; log as crash before diagnosis.","timestamp":1780356934871,"segment":0,"confidence":null,"packetFingerprint":"c3845ce12b7452161513667ad1dbf0504c56fd6e2606228b94c3505532d148d8","promotion":{"label":"blocked","reasons":["Crash evidence is retained without sentinel metrics and is blocked from promotion."]},"asi":{"hypothesis":"Re-run the current huge-document full target after fast-typing/runtime changes to verify aggregate budget.","evidence":"packet 4 parsed react_huge_doc_full_max_budget_ratio=1.06 but benchmark exited 1; failure_count=2; browser trace command failed while later overlay metrics emitted; burst_to_paint_p95_ms=139.6; type_to_paint_p95_ms=26.4; core_worst_p95_ratio=11.","rollback_reason":"No rollback: this packet did not apply a new change; benchmark failure blocks promotion.","next_action_hint":"Inspect tmp/slate-react-huge-document-full-benchmark-blocks-5000-iters-5-trace-iters-5-ops-10-full.json and browser-trace failure output; fix real browser-trace failures or benchmark harness before another packet.","lane":"react-huge-document-full","family":"browser-trace-failure","risk":"Metric parsed despite exit 1, so treat as crash, not improvement."},"benchmarkContract":{"command":"bash ./autoresearch.sh","checksCommand":"bash ./autoresearch.checks.sh","commandFile":"","envFile":"","surfaceHash":"168565aa676bf18c5ca56f781d6ff1ec47fe9bce2964216dbbfd959fa59389c2","files":[{"path":"autoresearch.sh","hash":"6d3d1a67661ec440d965ed38f0db82252434129a1b335e6ef961febd7c8e2bcd"},{"path":"autoresearch.checks.sh","hash":"9261aeaf280f5f66388ebfe2bea60e5dfb8d37ba8c69af4e608e7140f6fcdc20"},{"path":"package.json","hash":"4820e8774a8ee0fb005430a8af30b37f351200eb439460b10cd2ac47676252c2"}],"capturedAt":"2026-06-01T23:34:47.577Z"}}
{"run":5,"commit":"7eaafb9","metric":1.33,"metrics":{"react_huge_doc_type_to_paint_p95_ms":26.8,"react_huge_doc_dom_nodes_p95":20510,"react_huge_doc_heap_mb_p95":73.05,"react_huge_doc_long_task_max_p95_ms":0,"react_huge_doc_full_failure_count":2,"react_huge_doc_full_legacy_compare_worst_p95_ratio":0.77,"react_huge_doc_full_core_worst_p95_ratio":11,"react_huge_doc_full_type_to_paint_p95_ms":26.8,"react_huge_doc_full_burst_to_paint_p95_ms":139.4,"react_huge_doc_full_virtualized_type_to_paint_p95_ms":26.8,"react_huge_doc_full_dom_nodes_p95":20510,"react_huge_doc_full_long_task_max_p95_ms":0},"metricEligible":false,"status":"crash","evidenceStatus":"rejected","description":"Full huge-document packet still exits 1; partial-DOM promotion cold outlier remains after single-pass promotion and smaller segments.","timestamp":1780357448680,"segment":0,"confidence":null,"packetFingerprint":"d893446b62ccc8484f8d1f9500a09fddb2a099809d1106d64061cb771b0e6303","promotion":{"label":"blocked","reasons":["Crash evidence is retained without sentinel metrics and is blocked from promotion."]},"asi":{"hypothesis":"Synchronous segment-start selection plus smaller partial-DOM segments should remove the extra root-plan pass and reduce promotion mount cost.","evidence":"packet 5 parsed react_huge_doc_full_max_budget_ratio=1.33 but exited 1; overlay root-plan count is 1 and selector/text counts dropped, but partialDOMPromotionP95Ms is 66.36ms because the first fresh promotion sample is a cold outlier; type_to_paint=26.8ms and legacyCompareWorstP95Ratio=0.77 remain fine.","rollback_reason":"No rollback: runtime behavior tests passed and overlay steady samples improved, but full benchmark still blocks promotion.","next_action_hint":"Repair the overlay benchmark contract to separate cold first-promotion evidence from steady promotion latency, then rerun the full packet; do not keep until the benchmark exits 0 and checks pass.","lane":"react-huge-document-full","family":"partial-dom-promotion-measurement","risk":"Current p95 over 5 fresh mounts behaves like max and is dominated by one cold/noisy sample."},"benchmarkContract":{"command":"bash ./autoresearch.sh","checksCommand":"bash ./autoresearch.checks.sh","commandFile":"","envFile":"","surfaceHash":"168565aa676bf18c5ca56f781d6ff1ec47fe9bce2964216dbbfd959fa59389c2","files":[{"path":"autoresearch.sh","hash":"6d3d1a67661ec440d965ed38f0db82252434129a1b335e6ef961febd7c8e2bcd"},{"path":"autoresearch.checks.sh","hash":"9261aeaf280f5f66388ebfe2bea60e5dfb8d37ba8c69af4e608e7140f6fcdc20"},{"path":"package.json","hash":"4820e8774a8ee0fb005430a8af30b37f351200eb439460b10cd2ac47676252c2"}],"capturedAt":"2026-06-01T23:43:11.581Z"}}
{"run":6,"commit":"7eaafb9","metric":1.27,"metrics":{"react_huge_doc_legacy_compare_worst_p95_ratio":0.58,"react_huge_doc_legacy_compare_worst_p95_delta_ms":-19.83,"react_huge_doc_type_to_paint_p95_ms":22.7,"react_huge_doc_dom_nodes_p95":303,"react_huge_doc_heap_mb_p95":16.31,"react_huge_doc_long_task_max_p95_ms":0,"react_huge_doc_full_failure_count":0,"react_huge_doc_full_legacy_compare_worst_p95_ratio":0.58,"react_huge_doc_full_core_worst_p95_ratio":14,"react_huge_doc_full_type_to_paint_p95_ms":25.6,"react_huge_doc_full_burst_to_paint_p95_ms":139.5,"react_huge_doc_full_virtualized_type_to_paint_p95_ms":22.7,"react_huge_doc_full_partial_dom_promotion_steady_p95_ms":33.3,"react_huge_doc_full_partial_dom_promotion_cold_p95_ms":60.3,"react_huge_doc_full_dom_nodes_p95":20510,"react_huge_doc_full_long_task_max_p95_ms":0},"metricEligible":false,"status":"checks_failed","evidenceStatus":"rejected","description":"Benchmark clears with partial-DOM steady promotion and synthetic beforeinput excluded by default, but checks failed on formatter-only issues.","timestamp":1780358872785,"segment":0,"confidence":null,"packetFingerprint":"87c604b0b6aa998eb359464eba0beb171b4dcc083b09445cbd29a4011a1eb684","promotion":{"label":"blocked","reasons":["Correctness checks failed; packet evidence is blocked from promotion."]},"asi":{"hypothesis":"Excluding invalid synthetic beforeinput loops and fixing stale benchmark build inputs should let the huge-document packet measure the real V2 typing paths.","evidence":"packet 6 benchmark exit 0, react_huge_doc_full_max_budget_ratio=1.27, failure_count=0, virtualized_type_to_paint=22.7ms, partial_dom_promotion_steady=33.3ms, legacy_compare_worst_p95_ratio=0.58; checks failed only on formatter output in dom-strategy-and-scroll.tsx and repo-compare.mjs.","rollback_reason":"No rollback; benchmark and focused behavior tests passed, but packet cannot be promoted until formatting checks pass.","next_action_hint":"Apply formatter fixes, rerun focused tests and the full AR packet, then log as keep only if checks pass.","lane":"react-huge-document-full","family":"benchmark-contract-repair","risk":"Synthetic beforeinput remains opt-in because raw repeated beforeinput without real browser input pairing is not a valid typing oracle."},"benchmarkContract":{"command":"bash ./autoresearch.sh","checksCommand":"bash ./autoresearch.checks.sh","commandFile":"","envFile":"","surfaceHash":"168565aa676bf18c5ca56f781d6ff1ec47fe9bce2964216dbbfd959fa59389c2","files":[{"path":"autoresearch.sh","hash":"6d3d1a67661ec440d965ed38f0db82252434129a1b335e6ef961febd7c8e2bcd"},{"path":"autoresearch.checks.sh","hash":"9261aeaf280f5f66388ebfe2bea60e5dfb8d37ba8c69af4e608e7140f6fcdc20"},{"path":"package.json","hash":"4820e8774a8ee0fb005430a8af30b37f351200eb439460b10cd2ac47676252c2"}],"capturedAt":"2026-06-02T00:07:15.324Z"}}
{"run":7,"commit":"","metric":1.08,"metrics":{"react_huge_doc_legacy_compare_worst_p95_ratio":0.48,"react_huge_doc_legacy_compare_worst_p95_delta_ms":-29.71,"react_huge_doc_type_to_paint_p95_ms":23.3,"react_huge_doc_dom_nodes_p95":303,"react_huge_doc_heap_mb_p95":16.31,"react_huge_doc_long_task_max_p95_ms":0,"react_huge_doc_full_failure_count":0,"react_huge_doc_full_legacy_compare_worst_p95_ratio":0.48,"react_huge_doc_full_core_worst_p95_ratio":7,"react_huge_doc_full_type_to_paint_p95_ms":25.6,"react_huge_doc_full_burst_to_paint_p95_ms":141.7,"react_huge_doc_full_virtualized_type_to_paint_p95_ms":23.3,"react_huge_doc_full_partial_dom_promotion_steady_p95_ms":49.4,"react_huge_doc_full_partial_dom_promotion_cold_p95_ms":34.18,"react_huge_doc_full_dom_nodes_p95":20510,"react_huge_doc_full_long_task_max_p95_ms":0},"metricEligible":false,"status":"measure","evidenceStatus":"accepted","description":"Accepted measurement for huge-document benchmark contract repair and partial-DOM promotion improvements; packet passes benchmark and checks without committing.","timestamp":1780359137030,"segment":0,"confidence":null,"packetFingerprint":"98f7eb5dec5ecb5a04f26259ac5b265ce92868e26c445535669b27c8583beeb0","promotion":{"label":"measurement","reasons":["Logged as measure; metric evidence is trend-only and not finalizer evidence."]},"asi":{"hypothesis":"Synchronous partial-DOM promotion selection, smaller promotion islands, steady/cold promotion reporting, and removing invalid synthetic beforeinput from the default compare path should make the huge-document virtualized target measurable and faster without editor behavior regression.","evidence":"packet 7 passed benchmark and checks: react_huge_doc_full_max_budget_ratio=1.08 vs baseline 1.82; failure_count=0; virtualized_type_to_paint=23.3ms; type_to_paint=25.6ms; burst_to_paint=141.7ms; steady partial-DOM promotion=49.4ms; cold promotion=34.18ms; legacy compare worst p95 ratio=0.48; checks ran lint, typecheck, Bun tests, and Slate React Vitest green.","rollback_reason":"No rollback; packet improved the primary metric and correctness checks passed. Logged as measure rather than keep because keep requires an autoresearch commit and no commit was requested.","next_action_hint":"Next useful lane is reducing the remaining core_worst_p95_ratio=7 and dom_nodes_p95=20510, or splitting the full DOM-node budget by surface so virtualized/staged wins are not masked by full-surface DOM counts.","lane":"react-huge-document-full","family":"benchmark-contract-repair","risk":"Primary is still above 1.0 because core structural ratio and full-surface DOM node budget remain the limiting rows; do not claim the whole target is done."},"benchmarkContract":{"command":"bash ./autoresearch.sh","checksCommand":"bash ./autoresearch.checks.sh","commandFile":"","envFile":"","surfaceHash":"168565aa676bf18c5ca56f781d6ff1ec47fe9bce2964216dbbfd959fa59389c2","files":[{"path":"autoresearch.sh","hash":"6d3d1a67661ec440d965ed38f0db82252434129a1b335e6ef961febd7c8e2bcd"},{"path":"autoresearch.checks.sh","hash":"9261aeaf280f5f66388ebfe2bea60e5dfb8d37ba8c69af4e608e7140f6fcdc20"},{"path":"package.json","hash":"4820e8774a8ee0fb005430a8af30b37f351200eb439460b10cd2ac47676252c2"}],"capturedAt":"2026-06-02T00:11:41.558Z"}}
{"run":8,"commit":"","metric":1.02,"metrics":{"react_huge_doc_legacy_compare_worst_p95_ratio":0.42,"react_huge_doc_legacy_compare_worst_p95_delta_ms":-23.43,"react_huge_doc_type_to_paint_p95_ms":18.8,"react_huge_doc_dom_nodes_p95":303,"react_huge_doc_heap_mb_p95":16.31,"react_huge_doc_long_task_max_p95_ms":0,"react_huge_doc_full_failure_count":0,"react_huge_doc_full_legacy_compare_worst_p95_ratio":0.42,"react_huge_doc_full_core_worst_p95_ratio":12,"react_huge_doc_full_core_worst_budget_ratio":1.02,"react_huge_doc_full_type_to_paint_p95_ms":25,"react_huge_doc_full_burst_to_paint_p95_ms":139.3,"react_huge_doc_full_virtualized_type_to_paint_p95_ms":18.8,"react_huge_doc_full_partial_dom_promotion_steady_p95_ms":42.5,"react_huge_doc_full_partial_dom_promotion_cold_p95_ms":59.4,"react_huge_doc_full_dom_nodes_p95":20510,"react_huge_doc_full_browser_dom_nodes_p95":20510,"react_huge_doc_full_virtualized_dom_nodes_p95":303,"react_huge_doc_full_long_task_max_p95_ms":0},"metricEligible":false,"status":"measure","evidenceStatus":"accepted","description":"Accepted packet 8 measurement for huge-document full benchmark contract cleanup; benchmark and checks pass without committing.","timestamp":1780359693354,"segment":0,"confidence":null,"packetFingerprint":"6d21ab12f8df0c2fa309c745f853534fac50d4afafcbbfed84ef1cff570eee96","promotion":{"label":"measurement","reasons":["Logged as measure; metric evidence is trend-only and not finalizer evidence."]},"asi":{"hypothesis":"Separate budget-driving core evidence from diagnostic legacy ratios and split DOM-node metrics by full/browser versus virtualized surfaces so the huge-document target optimizes the real runtime constraint.","evidence":"packet 8 passed benchmark and checks: max_budget_ratio=1.02 vs packet 7 1.08 and baseline 1.82; failure_count=0; core_worst_budget_ratio=1.02; legacy_compare_worst_p95_ratio=0.42; virtualized_type_to_paint=18.8ms; type_to_paint=25ms; partial_dom_promotion_steady=42.5ms; browser_dom_nodes=20510 while virtualized_dom_nodes=303.","rollback_reason":"No rollback; checks passed. Logged as measure rather than keep because keep would invoke commit mechanics and no commit was requested.","next_action_hint":"The remaining primary miss is only core_worst_budget_ratio=1.02, likely noise at the current threshold. Repeat once or promote a gate with the cleaned budget metrics before deeper runtime edits.","lane":"react-huge-document-full","family":"benchmark-contract-cleanup","risk":"core_worst_p95_ratio remains high because selectAll is a sub-millisecond legacy-ratio diagnostic; do not treat it as a user-visible perf blocker."},"benchmarkContract":{"command":"bash ./autoresearch.sh","checksCommand":"bash ./autoresearch.checks.sh","commandFile":"","envFile":"","surfaceHash":"168565aa676bf18c5ca56f781d6ff1ec47fe9bce2964216dbbfd959fa59389c2","files":[{"path":"autoresearch.sh","hash":"6d3d1a67661ec440d965ed38f0db82252434129a1b335e6ef961febd7c8e2bcd"},{"path":"autoresearch.checks.sh","hash":"9261aeaf280f5f66388ebfe2bea60e5dfb8d37ba8c69af4e608e7140f6fcdc20"},{"path":"package.json","hash":"4820e8774a8ee0fb005430a8af30b37f351200eb439460b10cd2ac47676252c2"}],"capturedAt":"2026-06-02T00:21:07.167Z"}}
{"run":9,"commit":"","metric":1.22,"metrics":{"react_huge_doc_legacy_compare_worst_p95_ratio":0.66,"react_huge_doc_legacy_compare_worst_p95_delta_ms":-15.5,"react_huge_doc_type_to_paint_p95_ms":21,"react_huge_doc_dom_nodes_p95":303,"react_huge_doc_heap_mb_p95":16.31,"react_huge_doc_long_task_max_p95_ms":0,"react_huge_doc_full_failure_count":0,"react_huge_doc_full_legacy_compare_worst_p95_ratio":0.66,"react_huge_doc_full_core_worst_p95_ratio":16,"react_huge_doc_full_core_worst_budget_ratio":1.22,"react_huge_doc_full_type_to_paint_p95_ms":26,"react_huge_doc_full_burst_to_paint_p95_ms":139.6,"react_huge_doc_full_virtualized_type_to_paint_p95_ms":21,"react_huge_doc_full_partial_dom_promotion_steady_p95_ms":24.27,"react_huge_doc_full_partial_dom_promotion_cold_p95_ms":59.62,"react_huge_doc_full_dom_nodes_p95":20510,"react_huge_doc_full_browser_dom_nodes_p95":20510,"react_huge_doc_full_virtualized_dom_nodes_p95":303,"react_huge_doc_full_long_task_max_p95_ms":0},"metricEligible":false,"status":"measure","evidenceStatus":"accepted","description":"Accepted packet 9 repeat measurement; checks pass but the cleaned core budget ratio remains noisy and worse than packet 8.","timestamp":1780359933450,"segment":0,"confidence":null,"packetFingerprint":"d52a48a80af0bff4c6f3de6f6eb72a3fb5eb663abe65c297ac13e2751a8cbff1","promotion":{"label":"measurement","reasons":["Logged as measure; metric evidence is trend-only and not finalizer evidence."]},"asi":{"hypothesis":"Repeat the cleaned huge-document full benchmark once to determine whether packet 8 core_worst_budget_ratio=1.02 is stable enough for gate promotion.","evidence":"packet 9 passed benchmark and checks but worsened: max_budget_ratio=1.22, core_worst_budget_ratio=1.22, failure_count=0, virtualized_type_to_paint=21ms, type_to_paint=26ms, partial_dom_promotion_steady=24.27ms, browser_dom_nodes=20510, virtualized_dom_nodes=303. Runtime-facing React/virtualized lanes remain fine; the primary miss is core budget variance.","rollback_reason":"No rollback; repeat packet applied no new product change. Logged as measure because it is diagnostic repeat evidence, not a keep/commit candidate.","next_action_hint":"Do not patch React virtualization from this packet. Inspect the core huge-document budget row artifact and either make the core budget metric robust to sub-ms/noisy lanes or run a targeted core scout before another full packet.","lane":"react-huge-document-full","family":"repeat-validation","risk":"The aggregate primary is now dominated by core microbenchmark variance, while the user-visible browser and virtualized lanes remain comfortably under budget."},"benchmarkContract":{"command":"bash ./autoresearch.sh","checksCommand":"bash ./autoresearch.checks.sh","commandFile":"","envFile":"","surfaceHash":"168565aa676bf18c5ca56f781d6ff1ec47fe9bce2964216dbbfd959fa59389c2","files":[{"path":"autoresearch.sh","hash":"6d3d1a67661ec440d965ed38f0db82252434129a1b335e6ef961febd7c8e2bcd"},{"path":"autoresearch.checks.sh","hash":"9261aeaf280f5f66388ebfe2bea60e5dfb8d37ba8c69af4e608e7140f6fcdc20"},{"path":"package.json","hash":"4820e8774a8ee0fb005430a8af30b37f351200eb439460b10cd2ac47676252c2"}],"capturedAt":"2026-06-02T00:25:17.381Z"}}
{"run":10,"commit":"","metric":0.49,"metrics":{"react_huge_doc_legacy_compare_worst_p95_ratio":0.38,"react_huge_doc_legacy_compare_worst_p95_delta_ms":-56.29,"react_huge_doc_type_to_paint_p95_ms":23.4,"react_huge_doc_dom_nodes_p95":303,"react_huge_doc_heap_mb_p95":16.31,"react_huge_doc_long_task_max_p95_ms":0,"react_huge_doc_full_failure_count":0,"react_huge_doc_full_legacy_compare_worst_p95_ratio":0.38,"react_huge_doc_full_core_worst_p95_ratio":19,"react_huge_doc_full_core_worst_budget_ratio":0.28,"react_huge_doc_full_type_to_paint_p95_ms":26.4,"react_huge_doc_full_burst_to_paint_p95_ms":139.3,"react_huge_doc_full_virtualized_type_to_paint_p95_ms":23.4,"react_huge_doc_full_partial_dom_promotion_steady_p95_ms":24.7,"react_huge_doc_full_partial_dom_promotion_cold_p95_ms":58.85,"react_huge_doc_full_dom_nodes_p95":20510,"react_huge_doc_full_browser_dom_nodes_p95":20510,"react_huge_doc_full_virtualized_dom_nodes_p95":303,"react_huge_doc_full_long_task_max_p95_ms":0},"metricEligible":false,"status":"measure","evidenceStatus":"accepted","description":"Accept full huge-doc benchmark contract repair; primary ratio under budget with checks green.","timestamp":1780360443559,"segment":0,"confidence":null,"packetFingerprint":"18bca139ba3c081ad71ac4262ff4c3402822ac9f439f1886604ecd03066f6792","promotion":{"label":"measurement","reasons":["Logged as measure; metric evidence is trend-only and not finalizer evidence."]},"asi":{"hypothesis":"Gate tiny steady core edit lanes on robust p75 budget rows while preserving raw p95 diagnostics, so full huge-doc performance is judged by product-budget signal instead of scheduler noise.","evidence":"packet-10: react_huge_doc_full_max_budget_ratio=0.49; core_worst_budget_ratio=0.28; virtualized_type_to_paint=23.4ms; checks passed; raw core p95 ratio remains diagnostic at 19.","rollback_reason":"Revert if raw p95 diagnostics reveal real sustained over-budget work or if browser/editor behavior checks regress.","next_action_hint":"Run a repeat packet or move to the next concrete runtime bottleneck; do not chase core p95 outliers unless over-budget sample counts become sustained.","lane":"benchmark-contract","family":"huge-doc-full-metric-contract","risk":"measurement-policy"},"benchmarkContract":{"command":"bash ./autoresearch.sh","checksCommand":"bash ./autoresearch.checks.sh","commandFile":"","envFile":"","surfaceHash":"168565aa676bf18c5ca56f781d6ff1ec47fe9bce2964216dbbfd959fa59389c2","files":[{"path":"autoresearch.sh","hash":"6d3d1a67661ec440d965ed38f0db82252434129a1b335e6ef961febd7c8e2bcd"},{"path":"autoresearch.checks.sh","hash":"9261aeaf280f5f66388ebfe2bea60e5dfb8d37ba8c69af4e608e7140f6fcdc20"},{"path":"package.json","hash":"4820e8774a8ee0fb005430a8af30b37f351200eb439460b10cd2ac47676252c2"}],"capturedAt":"2026-06-02T00:33:40.982Z"}}
{"run":11,"commit":"7eaafb9","metric":0.92,"metrics":{"react_huge_doc_legacy_compare_worst_p95_ratio":0.61,"react_huge_doc_legacy_compare_worst_p95_delta_ms":-18.01,"react_huge_doc_type_to_paint_p95_ms":26,"react_huge_doc_burst_to_paint_per_op_p95_ms":14.78,"react_huge_doc_dom_nodes_p95":20510,"react_huge_doc_heap_mb_p95":77.63,"react_huge_doc_long_task_max_p95_ms":0,"react_huge_doc_full_failure_count":1,"react_huge_doc_full_legacy_compare_worst_p95_ratio":0.61,"react_huge_doc_full_core_worst_p95_ratio":17,"react_huge_doc_full_core_worst_budget_ratio":0.26,"react_huge_doc_full_type_to_paint_p95_ms":26,"react_huge_doc_full_burst_to_paint_p95_ms":147.8,"react_huge_doc_full_burst_to_paint_per_op_p95_ms":14.78,"react_huge_doc_full_virtualized_type_to_paint_p95_ms":26,"react_huge_doc_full_partial_dom_promotion_steady_p95_ms":33.55,"react_huge_doc_full_partial_dom_promotion_cold_p95_ms":60.05,"react_huge_doc_full_dom_nodes_p95":20510,"react_huge_doc_full_browser_dom_nodes_p95":20510,"react_huge_doc_full_virtualized_dom_nodes_p95":20510,"react_huge_doc_full_long_task_max_p95_ms":0},"metricEligible":false,"status":"crash","evidenceStatus":"rejected","description":"Full huge-doc packet exposed stale browser-trace artifact reuse after virtualized substep failed.","timestamp":1780360878571,"segment":0,"confidence":null,"packetFingerprint":"b03cb1f0ade41bdaad770ee263bb10135c2bc121540769dee051502cde048af1","promotion":{"label":"blocked","reasons":["Crash evidence is retained without sentinel metrics and is blocked from promotion."]},"asi":{"hypothesis":"Add normalized burst typing to the full huge-doc score.","evidence":"packet-11 exited 1; normalized burst metric emitted and passed at 14.78ms/op, but virtualized substep failed and full wrapper summarized the stale default/staged latest artifact as virtualized, reporting virtualizedDomNodes=20510.","rollback_reason":"Current wrapper reads tmp/slate-react-huge-document-browser-trace-benchmark.json for both browser trace substeps; stale latest artifacts make failed substeps look like real virtualized DOM regressions.","next_action_hint":"Fix full wrapper to read browser trace run-specific artifact paths per surface, then rerun focused tests and the full AR packet.","lane":"avoid","family":"huge-doc-full-artifact-contract","risk":"benchmark-integrity"},"benchmarkContract":{"command":"bash ./autoresearch.sh","checksCommand":"bash ./autoresearch.checks.sh","commandFile":"","envFile":"","surfaceHash":"168565aa676bf18c5ca56f781d6ff1ec47fe9bce2964216dbbfd959fa59389c2","files":[{"path":"autoresearch.sh","hash":"6d3d1a67661ec440d965ed38f0db82252434129a1b335e6ef961febd7c8e2bcd"},{"path":"autoresearch.checks.sh","hash":"9261aeaf280f5f66388ebfe2bea60e5dfb8d37ba8c69af4e608e7140f6fcdc20"},{"path":"package.json","hash":"4820e8774a8ee0fb005430a8af30b37f351200eb439460b10cd2ac47676252c2"}],"capturedAt":"2026-06-02T00:39:56.574Z"}}
{"run":12,"commit":"","metric":0.87,"metrics":{"react_huge_doc_legacy_compare_worst_p95_ratio":0.36,"react_huge_doc_legacy_compare_worst_p95_delta_ms":-27.12,"react_huge_doc_type_to_paint_p95_ms":23,"react_huge_doc_burst_to_paint_per_op_p95_ms":9.95,"react_huge_doc_dom_nodes_p95":303,"react_huge_doc_heap_mb_p95":16.31,"react_huge_doc_long_task_max_p95_ms":0,"react_huge_doc_full_failure_count":0,"react_huge_doc_full_legacy_compare_worst_p95_ratio":0.36,"react_huge_doc_full_core_worst_p95_ratio":166.5,"react_huge_doc_full_core_worst_budget_ratio":0.29,"react_huge_doc_full_type_to_paint_p95_ms":26.1,"react_huge_doc_full_burst_to_paint_p95_ms":139.4,"react_huge_doc_full_burst_to_paint_per_op_p95_ms":13.94,"react_huge_doc_full_virtualized_type_to_paint_p95_ms":23,"react_huge_doc_full_partial_dom_promotion_steady_p95_ms":39.74,"react_huge_doc_full_partial_dom_promotion_cold_p95_ms":55.14,"react_huge_doc_full_dom_nodes_p95":20510,"react_huge_doc_full_browser_dom_nodes_p95":20510,"react_huge_doc_full_virtualized_dom_nodes_p95":303,"react_huge_doc_full_long_task_max_p95_ms":0},"metricEligible":false,"status":"measure","evidenceStatus":"accepted","description":"Accept full huge-doc benchmark contract repair with normalized burst scoring and stale-artifact protection.","timestamp":1780361129132,"segment":0,"confidence":null,"packetFingerprint":"a5ae4a040208a1172b44af5f4bf4a3514bac045bbe13953628ca583e94ca6c56","promotion":{"label":"measurement","reasons":["Logged as measure; metric evidence is trend-only and not finalizer evidence."]},"asi":{"hypothesis":"Normalize burst typing per operation and make full huge-doc browser trace substeps read fresh run-specific artifacts, so the primary score reflects real per-character burst latency and failed substeps cannot reuse stale latest artifacts.","evidence":"packet-12: react_huge_doc_full_max_budget_ratio=0.87; failure_count=0; burst_per_op=13.94ms under 16ms budget; virtualized_type_to_paint=23ms; virtualized_dom_nodes=303; checks passed.","rollback_reason":"Revert if per-op burst hides sustained long tasks or if run-specific artifact paths diverge from the browser trace script naming contract.","next_action_hint":"Next bottleneck is partial DOM promotion steady p95=39.74ms and cold p95=55.14ms; core raw p95 remains diagnostic noise while core budget ratio is 0.29.","lane":"avoid","family":"huge-doc-full-artifact-contract","risk":"benchmark-integrity"},"benchmarkContract":{"command":"bash ./autoresearch.sh","checksCommand":"bash ./autoresearch.checks.sh","commandFile":"","envFile":"","surfaceHash":"168565aa676bf18c5ca56f781d6ff1ec47fe9bce2964216dbbfd959fa59389c2","files":[{"path":"autoresearch.sh","hash":"6d3d1a67661ec440d965ed38f0db82252434129a1b335e6ef961febd7c8e2bcd"},{"path":"autoresearch.checks.sh","hash":"9261aeaf280f5f66388ebfe2bea60e5dfb8d37ba8c69af4e608e7140f6fcdc20"},{"path":"package.json","hash":"4820e8774a8ee0fb005430a8af30b37f351200eb439460b10cd2ac47676252c2"}],"capturedAt":"2026-06-02T00:45:17.640Z"}}
{"type":"config","name":"react-huge-document-full","metricName":"react_huge_doc_full_max_budget_ratio","metricUnit":"ratio","bestDirection":"lower","segmentReason":"Start a fresh segment while preserving history.","timestamp":"2026-06-02T15:19:46.464Z"}
{"type":"lane_result","timestamp":1780413794053,"segment":1,"lane":{"id":"read-only-scout","title":"Read-only scout","mode":"read_only_scout"},"result":{"status":"completed","summary":"Read-only scout: latest historical budget rows show browserTraceBurstToPaintPerOpP95Ms is the current max row at 13.94ms/16ms ratio 0.87; partialDOMPromotionSteadyP95Ms is next at 39.74ms/50ms ratio 0.79; virtualized type-to-paint and DOM nodes are comfortably under budget.","recommendation":"Run a fresh segment-1 baseline before editing. If burst per-op remains the max row, target browser trace burst typing; if it drops and partial DOM promotion becomes max, target partial-DOM promotion.","evidenceAccepted":true,"command":"","timeBudgetSeconds":300,"isolation":{"mode":"read_only_scout","worktree":"","writeScope":[]},"commandResult":null}}
{"run":13,"commit":"","metric":0.82,"metrics":{"react_huge_doc_legacy_compare_worst_p95_ratio":0.36,"react_huge_doc_legacy_compare_worst_p95_delta_ms":-55.53,"react_huge_doc_type_to_paint_p95_ms":31.1,"react_huge_doc_burst_to_paint_per_op_p95_ms":3.72,"react_huge_doc_dom_nodes_p95":303,"react_huge_doc_heap_mb_p95":16.31,"react_huge_doc_long_task_max_p95_ms":0,"react_huge_doc_full_failure_count":0,"react_huge_doc_full_legacy_compare_worst_p95_ratio":0.36,"react_huge_doc_full_core_worst_p95_ratio":10.5,"react_huge_doc_full_core_worst_budget_ratio":0.4,"react_huge_doc_full_type_to_paint_p95_ms":33.4,"react_huge_doc_full_burst_to_paint_p95_ms":131.4,"react_huge_doc_full_burst_to_paint_per_op_p95_ms":13.14,"react_huge_doc_full_virtualized_type_to_paint_p95_ms":31.1,"react_huge_doc_full_partial_dom_promotion_steady_p95_ms":23.73,"react_huge_doc_full_partial_dom_promotion_cold_p95_ms":32.81,"react_huge_doc_full_dom_nodes_p95":20510,"react_huge_doc_full_browser_dom_nodes_p95":20510,"react_huge_doc_full_virtualized_dom_nodes_p95":303,"react_huge_doc_full_long_task_max_p95_ms":0},"metricEligible":false,"status":"measure","evidenceStatus":"accepted","description":"Accepted segment-1 baseline for react-huge-document-full; benchmark and checks pass with max budget ratio under budget.","timestamp":1780413987168,"segment":1,"confidence":null,"packetFingerprint":"396aa1dc782c80bf929aaca5fb25685e260c160650526b00aa501736936bba16","promotion":{"label":"measurement","reasons":["Logged as measure; metric evidence is trend-only and not finalizer evidence."]},"asi":{"hypothesis":"A fresh segment-1 baseline should confirm the current huge-document full budget after segment transition before any further perf edit.","evidence":"packet-13: react_huge_doc_full_max_budget_ratio=0.82; failure_count=0; checks passed; burst_per_op=13.14ms under 16ms; virtualized_type_to_paint=31.1ms; partial_dom_promotion_steady=23.73ms; partial_dom_promotion_cold=32.81ms; virtualized_dom_nodes=303.","lane":"react-huge-document-full","family":"segment-1-baseline","next_action_hint":"All budget rows are under target. If optimizing further, first scout browserTraceBurstToPaintPerOpP95Ms because it is the current max row; otherwise run promotion-readiness/repeat before claiming final perf."},"benchmarkContract":{"command":"bash ./autoresearch.sh","checksCommand":"bash ./autoresearch.checks.sh","commandFile":"","envFile":"","surfaceHash":"168565aa676bf18c5ca56f781d6ff1ec47fe9bce2964216dbbfd959fa59389c2","files":[{"path":"autoresearch.sh","hash":"6d3d1a67661ec440d965ed38f0db82252434129a1b335e6ef961febd7c8e2bcd"},{"path":"autoresearch.checks.sh","hash":"9261aeaf280f5f66388ebfe2bea60e5dfb8d37ba8c69af4e608e7140f6fcdc20"},{"path":"package.json","hash":"4820e8774a8ee0fb005430a8af30b37f351200eb439460b10cd2ac47676252c2"}],"capturedAt":"2026-06-02T15:26:07.085Z"}}
{"type":"lane_result","timestamp":1780414069276,"segment":1,"lane":{"id":"explore","title":"Explore","mode":"read_only_scout"},"result":{"status":"completed","summary":"Explore scout: the current max budget row is browserTraceBurstToPaintPerOpP95Ms=13.14ms/16ms, owned by stagedDomPresent/startBlock. defaultAuto is 3.94ms/op start and 3.44ms/op middle; virtualized aggregate is 3.72ms/op. No long tasks. Partial DOM promotion is no longer close to budget: steady=23.73ms/50ms, cold=32.81ms. The hot row is full-DOM staged diagnostic work, not virtualized product typing.","recommendation":"Do not patch yet. Next owner should be promotion-readiness or a repeat gate: one more segment-1 packet or promote-gate dry run to prove the under-budget result is stable before claiming final huge-document perf. Only optimize stagedDomPresent burst if repeat regresses near or above 16ms/op.","evidenceAccepted":true,"command":"","timeBudgetSeconds":300,"isolation":{"mode":"read_only_scout","worktree":"","writeScope":[]},"commandResult":null}}
{"type":"lane_result","timestamp":1780414141558,"segment":1,"lane":{"id":"benchmark-contract","title":"Benchmark contract","mode":"read_only_scout"},"result":{"status":"completed","summary":"Benchmark contract scout: autoresearch.sh runs react-huge-document-full with 5000 blocks, 5 iterations, 5 trace iterations, 10 type ops, and legacy repo ../../../slate. Captured benchmark output emits react_huge_doc_full_max_budget_ratio plus failure_count and secondary metrics; benchmark-lint parsed the primary metric cleanly. The full artifact has 15 finite budget rows, max row browserTraceBurstToPaintPerOpP95Ms=13.14ms/16ms ratio 0.82, invalidRows=0. Checks command is bun check and passed during packet 13.","recommendation":"Benchmark contract is OK. Next owner should be promotion-readiness: run a repeat or promote-gate dry run because research integrity still lacks holdout/repeat/promotion guard evidence. Do not patch runtime until repeat evidence breaks budget or identifies a real product-surface regression.","evidenceAccepted":true,"command":"","timeBudgetSeconds":300,"isolation":{"mode":"read_only_scout","worktree":"","writeScope":[]},"commandResult":null}}
{"type":"lane_result","timestamp":1780415612100,"segment":1,"lane":{"id":"implementation-candidate","title":"Implementation candidate","mode":"implementation"},"result":{"status":"completed","summary":"Implementation candidate: no source patch is justified. Segment-1 baseline is under budget at max_budget_ratio=0.82; the max row is stagedDomPresent/startBlock burst typing at 13.14ms/op against a 16ms budget; defaultAuto and virtualized lanes are much faster, partial DOM promotion is comfortably under budget, and checks passed. A runtime patch now would optimize a diagnostic full-DOM row without a failing product budget.","recommendation":"Do not edit product/runtime code. Move to promotion-readiness: run one repeat packet or promote-gate dry run to prove the under-budget result is stable. Reopen implementation only if repeat evidence exceeds budget or exposes a correctness regression.","evidenceAccepted":true,"command":"","timeBudgetSeconds":300,"isolation":{"mode":"implementation","worktree":"","writeScope":["autoresearch.jsonl","autoresearch.md"]},"commandResult":null}}
{"type":"lane_result","timestamp":1780415631101,"segment":1,"lane":{"id":"promotion-readiness","title":"Promotion readiness","mode":"read_only_scout"},"result":{"status":"completed","summary":"Promotion-readiness criterion: segment 1 has one accepted baseline measure at max_budget_ratio=0.82 with checks passing and no implementation candidate. Promotion is not strong enough from one packet alone; it needs repeat evidence or a promote-gate dry run before commit readiness is honest.","recommendation":"Run exactly one repeat packet with autoresearch next. If it remains under budget with failure_count=0 and checks pass, log it as accepted measure and stop at commit approval. If it regresses over budget, reopen implementation-candidate only for the new max row.","evidenceAccepted":true,"command":"","timeBudgetSeconds":300,"isolation":{"mode":"read_only_scout","worktree":"","writeScope":[]},"commandResult":null}}
{"run":14,"commit":"","metric":0.82,"metrics":{"react_huge_doc_legacy_compare_worst_p95_ratio":0.74,"react_huge_doc_legacy_compare_worst_p95_delta_ms":-22.33,"react_huge_doc_type_to_paint_p95_ms":30.9,"react_huge_doc_burst_to_paint_per_op_p95_ms":3.2,"react_huge_doc_dom_nodes_p95":303,"react_huge_doc_heap_mb_p95":16.31,"react_huge_doc_long_task_max_p95_ms":0,"react_huge_doc_full_failure_count":0,"react_huge_doc_full_legacy_compare_worst_p95_ratio":0.74,"react_huge_doc_full_core_worst_p95_ratio":13.5,"react_huge_doc_full_core_worst_budget_ratio":0.38,"react_huge_doc_full_type_to_paint_p95_ms":36.8,"react_huge_doc_full_burst_to_paint_p95_ms":131.3,"react_huge_doc_full_burst_to_paint_per_op_p95_ms":13.13,"react_huge_doc_full_virtualized_type_to_paint_p95_ms":30.9,"react_huge_doc_full_partial_dom_promotion_steady_p95_ms":25.89,"react_huge_doc_full_partial_dom_promotion_cold_p95_ms":58.38,"react_huge_doc_full_dom_nodes_p95":20510,"react_huge_doc_full_browser_dom_nodes_p95":20510,"react_huge_doc_full_virtualized_dom_nodes_p95":303,"react_huge_doc_full_long_task_max_p95_ms":0},"metricEligible":false,"status":"measure","evidenceStatus":"accepted","description":"Accepted segment-1 repeat for react-huge-document-full; benchmark and checks pass with stable max budget ratio under budget.","timestamp":1780415897999,"segment":1,"confidence":null,"packetFingerprint":"7c034e3dd1d71c8de1b13fba3c0e777143e0e5eb30e5e3ed8de33cc57d0da95a","promotion":{"label":"measurement","reasons":["Logged as measure; metric evidence is trend-only and not finalizer evidence."]},"asi":{"hypothesis":"Repeat packet should prove the segment-1 huge-document full baseline is stable enough for commit readiness without a runtime patch.","evidence":"packet-14: react_huge_doc_full_max_budget_ratio=0.82; failure_count=0; checks passed; burst_per_op=13.13ms under 16ms; virtualized_type_to_paint=30.9ms; partial_dom_promotion_steady=25.89ms; partial_dom_promotion_cold=58.38ms; virtualized_dom_nodes=303.","lane":"react-huge-document-full","family":"segment-1-repeat-readiness","next_action_hint":"Promotion-readiness repeat passed. No source patch is justified; next autonomous step is finalization/commit-readiness preview, then stop for commit approval."},"benchmarkContract":{"command":"bash ./autoresearch.sh","checksCommand":"bash ./autoresearch.checks.sh","commandFile":"","envFile":"","surfaceHash":"168565aa676bf18c5ca56f781d6ff1ec47fe9bce2964216dbbfd959fa59389c2","files":[{"path":"autoresearch.sh","hash":"6d3d1a67661ec440d965ed38f0db82252434129a1b335e6ef961febd7c8e2bcd"},{"path":"autoresearch.checks.sh","hash":"9261aeaf280f5f66388ebfe2bea60e5dfb8d37ba8c69af4e608e7140f6fcdc20"},{"path":"package.json","hash":"4820e8774a8ee0fb005430a8af30b37f351200eb439460b10cd2ac47676252c2"}],"capturedAt":"2026-06-02T15:56:48.411Z"}}