{"forums": [{"scenario": "a softening", "forum": "LwjUKEWAvt", "title": "SafetyChat: Learning to Generate Physical Safety Warnings in Instructional Assistants", "decision": "Reject", "comments": 12, "reviewers": [{"key": "DUDk", "rating": 6, "units": [{"i": 0, "phase": "init", "chg": null, "v": "negative", "o": "baselines_ablations", "t": "Experiment lacks completeness."}, {"i": 1, "phase": "init", "chg": null, "v": "negative", "o": "reproducibility", "t": "Cannot assess reproducibility."}, {"i": 2, "phase": "init", "chg": null, "v": "negative", "o": "empirical_scope", "t": "Generalization validity concern."}, {"i": 3, "phase": "init", "chg": null, "v": "negative", "o": "reproducibility", "t": "Insufficient documentation."}, {"i": 4, "phase": "post", "chg": "weakened", "v": "mixed", "o": "baselines_ablations", "t": "Concerns partially addressed."}, {"i": 5, "phase": "post", "chg": "weakened", "v": "positive", "o": "reproducibility", "t": "Reproducibility concern resolved."}]}, {"key": "RX83", "rating": 4, "units": [{"i": 0, "phase": "init", "chg": null, "v": "negative", "o": "empirical_scope", "t": "Limited generalization validity; alternative explanation for performance exists."}, {"i": 1, "phase": "init", "chg": null, "v": "negative", "o": "reproducibility", "t": "Construct validity concern."}, {"i": 2, "phase": "init", "chg": null, "v": "negative", "o": "novelty", "t": "Incomplete experimental design."}, {"i": 3, "phase": "init", "chg": null, "v": "negative", "o": "stats_metrics", "t": "Validity-of-evaluation concern."}, {"i": 4, "phase": "init", "chg": null, "v": "negative", "o": "clarity", "t": "Minor presentation/access issue."}, {"i": 5, "phase": "post", "chg": "weakened", "v": "negative", "o": "empirical_scope", "t": "Generalization concern not fully resolved."}, {"i": 6, "phase": "post", "chg": "weakened", "v": "positive", "o": "baselines_ablations", "t": "Baseline concern resolved."}, {"i": 7, "phase": "post", "chg": "weakened", "v": "mixed", "o": "method_design", "t": "Documentation improved, but construct validity concern persists."}]}, {"key": "6je4", "rating": 4, "units": [{"i": 0, "phase": "init", "chg": null, "v": "negative", "o": "stats_metrics", "t": "Critical validity flaw due to circular evaluation."}, {"i": 1, "phase": "init", "chg": null, "v": "negative", "o": "empirical_scope", "t": "Incomplete scope of evaluation."}, {"i": 2, "phase": "init", "chg": null, "v": "negative", "o": "empirical_scope", "t": "Significant overclaim in framing."}, {"i": 3, "phase": "init", "chg": null, "v": "negative", "o": "empirical_scope", "t": "Ambiguity in what was learned."}, {"i": 4, "phase": "post", "chg": "weakened", "v": "mixed", "o": "stats_metrics", "t": "Paradox partially mitigated but not fully resolved."}, {"i": 5, "phase": "post", "chg": null, "v": "negative", "o": "clarity", "t": "Multimodal concern not resolved."}, {"i": 6, "phase": "post", "chg": "weakened", "v": "positive", "o": "related_work", "t": "Framing concern resolved."}, {"i": 7, "phase": "post", "chg": "weakened", "v": "mixed", "o": "empirical_scope", "t": "Generalization concern partially addressed but scope remains limited."}]}, {"key": "Lqwh", "rating": 2, "units": [{"i": 0, "phase": "init", "chg": null, "v": "negative", "o": "related_work", "t": "Insufficient novelty argument; the burden of differentiation is unmet."}, {"i": 1, "phase": "init", "chg": null, "v": "negative", "o": "clarity", "t": "Incomplete representation of the task domain."}, {"i": 2, "phase": "init", "chg": null, "v": "negative", "o": "empirical_scope", "t": "Methodological concern regarding bias."}, {"i": 3, "phase": "init", "chg": null, "v": "negative", "o": "baselines_ablations", "t": "Incomplete experimental comparison."}, {"i": 4, "phase": "init", "chg": null, "v": "negative", "o": "theory", "t": "Insufficient algorithmic novelty."}, {"i": 5, "phase": "init", "chg": null, "v": "negative", "o": "clarity", "t": "Presentation flaws."}, {"i": 6, "phase": "init", "chg": null, "v": "negative", "o": "theory", "t": "Potential training instability due to data mismatch."}, {"i": 7, "phase": "post", "chg": "strengthened", "v": "negative", "o": "clarity", "t": "Multimodal concern not fully resolved."}, {"i": 8, "phase": "post", "chg": "weakened", "v": "negative", "o": "stats_metrics", "t": "Bias concern partially addressed but not fully resolved."}, {"i": 9, "phase": "post", "chg": "weakened", "v": "positive", "o": "baselines_ablations", "t": "Baseline concern resolved."}, {"i": 10, "phase": "post", "chg": null, "v": "mixed", "o": "theory", "t": "Scope clarified, novelty expectation met by dataset value."}]}]}, {"scenario": "a softening", "forum": "hZnibTOke7", "title": "Self-Speculative Decoding Accelerates Lossless Inference in Any-Order and Any-Subset Autoregress", "decision": "Accept (Poster)", "comments": 8, "reviewers": [{"key": "ummp", "rating": 8, "units": [{"i": 0, "phase": "init", "chg": null, "v": "positive", "o": "theory", "t": "The paper offers a significant qualitative contribution through its theoretical soundness."}, {"i": 1, "phase": "init", "chg": null, "v": "positive", "o": "theory", "t": "The empirical results support the theoretical framework."}, {"i": 2, "phase": "init", "chg": null, "v": "positive", "o": "baselines_ablations", "t": "The authors demonstrate high research integrity."}, {"i": 3, "phase": "init", "chg": null, "v": "negative", "o": "baselines_ablations", "t": "The empirical story is incomplete without comparative benchmarks."}, {"i": 4, "phase": "init", "chg": null, "v": "negative", "o": "robustness_sensitivity", "t": "Methodological choices lack sufficient justification."}, {"i": 5, "phase": "init", "chg": null, "v": "negative", "o": "empirical_scope", "t": "Implementation details regarding attention patterns are ambiguous."}, {"i": 6, "phase": "init", "chg": null, "v": "negative", "o": "stats_metrics", "t": "Performance patterns require explanation regarding training dynamics."}, {"i": 7, "phase": "init", "chg": null, "v": "mixed", "o": "compute_cost", "t": "Speedup is limited but acceptable given guarantees."}, {"i": 8, "phase": "init", "chg": null, "v": "negative", "o": "method_design", "t": "Dependence on pretrained initialization needs clarification."}, {"i": 9, "phase": "post", "chg": "weakened", "v": "positive", "o": "baselines_ablations", "t": "The missing comparison weakness was substantially addressed."}, {"i": 10, "phase": "post", "chg": "weakened", "v": "positive", "o": "robustness_sensitivity", "t": "The hyperparameter concern was clarified and weakened."}, {"i": 11, "phase": "post", "chg": "weakened", "v": "positive", "o": "empirical_scope", "t": "The implementation ambiguity was resolved."}, {"i": 12, "phase": "post", "chg": "weakened", "v": "positive", "o": "stats_metrics", "t": "The counter-intuitive results were explained satisfactorily."}, {"i": 13, "phase": "post", "chg": "weakened", "v": "positive", "o": "method_design", "t": "The dependence on initialization was clarified and justified."}]}, {"key": "yrNB", "rating": 8, "units": [{"i": 0, "phase": "init", "chg": null, "v": "negative", "o": "stats_metrics", "t": "Practical value is limited by small acceptance length."}, {"i": 1, "phase": "init", "chg": null, "v": "positive", "o": "compute_cost", "t": "Resource constraints mitigate the severity of the acceptance length issue."}, {"i": 2, "phase": "init", "chg": null, "v": "negative", "o": "compute_cost", "t": "Comparison fairness is questionable."}, {"i": 3, "phase": "post", "chg": "weakened", "v": "mixed", "o": "compute_cost", "t": "The fairness concern was nuanced but not fully resolved; reviewer downgraded importance."}, {"i": 4, "phase": "post", "chg": "weakened", "v": "mixed", "o": "compute_cost", "t": "Speedup limitation remained but was discounted."}]}, {"key": "BarN", "rating": 4, "units": [{"i": 0, "phase": "init", "chg": null, "v": "negative", "o": "theory", "t": "The paper feels incomplete due to narrow focus."}, {"i": 1, "phase": "init", "chg": null, "v": "negative", "o": "stats_metrics", "t": "Empirical story is incomplete without accuracy analysis."}, {"i": 2, "phase": "init", "chg": null, "v": "negative", "o": "baselines_ablations", "t": "Uncertainty remains about trade-offs between speculation and quality."}, {"i": 3, "phase": "post", "chg": "weakened", "v": "positive", "o": "theory", "t": "The scope concern was addressed."}, {"i": 4, "phase": "post", "chg": "weakened", "v": "positive", "o": "stats_metrics", "t": "The accuracy concern was clarified."}]}]}, {"scenario": "a softening", "forum": "h8u0KWgg9C", "title": "Offline Equilibrium Finding in Extensive-form Games: Datasets, Methods, and Analysis", "decision": "Reject", "comments": 15, "reviewers": [{"key": "wSAG", "rating": 6, "units": [{"i": 0, "phase": "init", "chg": null, "v": "negative", "o": "theory", "t": "The theoretical analysis in Appendix D contains a gap regarding MB performance under unilateral coverage."}, {"i": 1, "phase": "init", "chg": null, "v": "negative", "o": "compute_cost", "t": "The experimental design exhibits a theoretical mismatch in algorithm usage."}, {"i": 2, "phase": "init", "chg": "weakened", "v": "negative", "o": "theory", "t": "The evaluation metric is ambiguously defined for the multiplayer context."}, {"i": 3, "phase": "init", "chg": null, "v": "negative", "o": "compute_cost", "t": "The problem is computationally hard, raising the bar for theoretical contributions."}, {"i": 4, "phase": "init", "chg": null, "v": "negative", "o": "clarity", "t": "Minor terminological imprecisions present."}, {"i": 5, "phase": "init", "chg": "weakened", "v": "negative", "o": "compute_cost", "t": "The relationship between coverage assumptions needs clarification."}, {"i": 6, "phase": "post", "chg": null, "v": "negative", "o": "robustness_sensitivity", "t": "The authors' theoretical gap can be bridged by pessimistic assumptions, which they did not address."}, {"i": 7, "phase": "post", "chg": null, "v": "negative", "o": "theory", "t": "The methodological inconsistency remains theoretically problematic."}]}, {"key": "yN9Z", "rating": 6, "units": [{"i": 0, "phase": "init", "chg": "weakened", "v": "negative", "o": "theory", "t": "The related work section lacks sufficient depth on Markov games."}, {"i": 1, "phase": "init", "chg": "weakened", "v": "negative", "o": "related_work", "t": "The related work section misses relevant applied research."}, {"i": 2, "phase": "init", "chg": "weakened", "v": "negative", "o": "reproducibility", "t": "The related work section lacks discussion on offline MARL."}]}, {"key": "nYzC", "rating": 4, "units": [{"i": 0, "phase": "init", "chg": null, "v": "negative", "o": "robustness_sensitivity", "t": "The mixing parameter estimation mechanism is theoretically weak."}, {"i": 1, "phase": "init", "chg": null, "v": "negative", "o": "theory", "t": "The theoretical guarantee for BOMB's superiority over BC is unsupported or incorrect."}, {"i": 2, "phase": "init", "chg": null, "v": "negative", "o": "theory", "t": "The claim of algorithm independence is insufficiently supported."}, {"i": 3, "phase": "init", "chg": "weakened", "v": "negative", "o": "clarity", "t": "The manuscript contains noticeable errors."}]}, {"key": "XxFT", "rating": 2, "units": [{"i": 0, "phase": "init", "chg": null, "v": "negative", "o": "theory", "t": "The theoretical assumptions are unrealistic for practical settings."}, {"i": 1, "phase": "init", "chg": null, "v": "negative", "o": "stats_metrics", "t": "The experimental validation does not sufficiently demonstrate scalability or practical utility."}, {"i": 2, "phase": "init", "chg": "weakened", "v": "negative", "o": "empirical_scope", "t": "The dataset generation approach conflicts with the stated problem motivation."}, {"i": 3, "phase": "init", "chg": null, "v": "negative", "o": "method_design", "t": "The novelty claim is inaccurate due to insufficient literature review."}, {"i": 4, "phase": "init", "chg": "weakened", "v": "negative", "o": "novelty", "t": "The core algorithmic contribution lacks novelty."}, {"i": 5, "phase": "init", "chg": null, "v": "negative", "o": "stats_metrics", "t": "The MB strategy has convergence issues in specific unseen-state scenarios."}, {"i": 6, "phase": "init", "chg": null, "v": "negative", "o": "related_work", "t": "The related work section is incomplete and outdated."}, {"i": 7, "phase": "init", "chg": null, "v": "negative", "o": "theory", "t": "The evaluation metric is flawed for non-unique equilibrium settings."}, {"i": 8, "phase": "init", "chg": null, "v": "negative", "o": "empirical_scope", "t": "The state encoding method is insufficiently explained."}, {"i": 9, "phase": "init", "chg": null, "v": "negative", "o": "compute_cost", "t": "Scalability is not adequately addressed."}, {"i": 10, "phase": "init", "chg": null, "v": "negative", "o": "stats_metrics", "t": "The comparative analysis is incomplete."}, {"i": 11, "phase": "init", "chg": null, "v": "negative", "o": "baselines_ablations", "t": "The experimental comparison set is incomplete."}, {"i": 12, "phase": "post", "chg": "strengthened", "v": "positive", "o": "method_design", "t": "The fundamental contribution is recognized as innovative."}, {"i": 13, "phase": "post", "chg": "weakened", "v": "positive", "o": "empirical_scope", "t": "The previous concern about dataset mixing is invalid."}, {"i": 14, "phase": "post", "chg": null, "v": "negative", "o": "compute_cost", "t": "The method lacks demonstrated practical utility."}, {"i": 15, "phase": "post", "chg": null, "v": "negative", "o": "theory", "t": "Convergence in the counterexample scenario remains unresolved."}]}]}, {"scenario": "an entrenchment", "forum": "zfVICPB5Sv", "title": "Silent Leaks: Implicit Knowledge Extraction Attack on RAG Systems", "decision": "Accept (Poster)", "comments": 25, "reviewers": [{"key": "iS3Y", "rating": 6, "units": [{"i": 0, "phase": "init", "chg": null, "v": "positive", "o": "empirical_scope", "t": "Positive evaluation of the problem formulation and approach novelty."}, {"i": 1, "phase": "init", "chg": null, "v": "positive", "o": "clarity", "t": "Positive evaluation of presentation quality."}, {"i": 2, "phase": "init", "chg": null, "v": "positive", "o": "method_design", "t": "Positive evaluation of empirical results against weak baselines."}, {"i": 3, "phase": "init", "chg": null, "v": "positive", "o": "reproducibility", "t": "Positive evaluation of research openness."}, {"i": 4, "phase": "init", "chg": "strengthened", "v": "negative", "o": "stats_metrics", "t": "Negative evaluation of defense realism and stealth validation."}, {"i": 5, "phase": "init", "chg": "strengthened", "v": "negative", "o": "empirical_scope", "t": "Skeptical evaluation of generalizability due to topic constraint."}, {"i": 6, "phase": "init", "chg": "strengthened", "v": "negative", "o": "stats_metrics", "t": "Negative evaluation of empirical support for functional equivalence claims."}, {"i": 7, "phase": "init", "chg": "strengthened", "v": "negative", "o": "compute_cost", "t": "Negative evaluation of economic feasibility transparency."}, {"i": 8, "phase": "init", "chg": "strengthened", "v": "positive", "o": "empirical_scope", "t": "Neutral inquiry into operational stability."}, {"i": 9, "phase": "init", "chg": null, "v": "positive", "o": "robustness_sensitivity", "t": "Neutral inquiry into evasion durability."}, {"i": 10, "phase": "post", "chg": "strengthened", "v": "positive", "o": "method_design", "t": "Satisfaction with the improvements made to the paper."}]}, {"key": "CR94", "rating": 4, "units": [{"i": 0, "phase": "init", "chg": null, "v": "positive", "o": "method_design", "t": "Positive evaluation of the experimental setup realism."}, {"i": 1, "phase": "init", "chg": null, "v": "positive", "o": "empirical_scope", "t": "Positive evaluation of comparative performance."}, {"i": 2, "phase": "init", "chg": null, "v": "positive", "o": "clarity", "t": "Positive evaluation of presentation."}, {"i": 3, "phase": "init", "chg": "strengthened", "v": "negative", "o": "empirical_scope", "t": "Negative evaluation of theoretical grounding."}, {"i": 4, "phase": "init", "chg": "strengthened", "v": "negative", "o": "baselines_ablations", "t": "Negative evaluation of experimental completeness."}, {"i": 5, "phase": "init", "chg": "strengthened", "v": "negative", "o": "stats_metrics", "t": "Negative evaluation of practical threat magnitude."}, {"i": 6, "phase": "init", "chg": "strengthened", "v": "negative", "o": "stats_metrics", "t": "Negative evaluation of measurement precision."}, {"i": 7, "phase": "init", "chg": "strengthened", "v": "negative", "o": "robustness_sensitivity", "t": "Negative evaluation of method stability."}, {"i": 8, "phase": "init", "chg": "strengthened", "v": "negative", "o": "baselines_ablations", "t": "Negative evaluation of component effectiveness."}, {"i": 9, "phase": "init", "chg": "strengthened", "v": "positive", "o": "empirical_scope", "t": "Neutral inquiry into metric behavior."}]}, {"key": "9Ln3", "rating": 4, "units": [{"i": 0, "phase": "init", "chg": null, "v": "positive", "o": "novelty", "t": "Positive evaluation of presentation and significance."}, {"i": 1, "phase": "init", "chg": "strengthened", "v": "negative", "o": "baselines_ablations", "t": "Negative evaluation of experimental fairness."}, {"i": 2, "phase": "init", "chg": "strengthened", "v": "negative", "o": "stats_metrics", "t": "Negative evaluation of metric design neutrality."}, {"i": 3, "phase": "init", "chg": "strengthened", "v": "negative", "o": "empirical_scope", "t": "Negative evaluation of external validity."}, {"i": 4, "phase": "init", "chg": "strengthened", "v": "negative", "o": "related_work", "t": "Negative evaluation of end-to-end automation."}, {"i": 5, "phase": "init", "chg": "strengthened", "v": "negative", "o": "robustness_sensitivity", "t": "Negative evaluation of defense utility."}, {"i": 6, "phase": "init", "chg": "strengthened", "v": "negative", "o": "novelty", "t": "Negative evaluation of claim substantiation."}, {"i": 7, "phase": "init", "chg": "strengthened", "v": "negative", "o": "baselines_ablations", "t": "Conditional negative pending evidence."}]}, {"key": "LmNX", "rating": 4, "units": [{"i": 0, "phase": "init", "chg": null, "v": "positive", "o": "empirical_scope", "t": "Positive evaluation of threat model precision."}, {"i": 1, "phase": "init", "chg": null, "v": "positive", "o": "stats_metrics", "t": "Positive evaluation of experimental scope."}, {"i": 2, "phase": "init", "chg": null, "v": "positive", "o": "reproducibility", "t": "Positive evaluation of transparency."}, {"i": 3, "phase": "init", "chg": "strengthened", "v": "negative", "o": "novelty", "t": "Negative evaluation of theoretical contribution."}, {"i": 4, "phase": "init", "chg": "strengthened", "v": "negative", "o": "empirical_scope", "t": "Negative evaluation of generalizability."}, {"i": 5, "phase": "init", "chg": "strengthened", "v": "negative", "o": "method_design", "t": "Negative evaluation of defense coverage."}, {"i": 6, "phase": "init", "chg": "strengthened", "v": "negative", "o": "stats_metrics", "t": "Negative evaluation of economic transparency."}]}]}, {"scenario": "an entrenchment", "forum": "7LoFonLZqs", "title": "Greater than the Sum of Its Parts:  Building Substructure into Protein Encoding Models", "decision": "Accept (Poster)", "comments": 18, "reviewers": [{"key": "w8Xw", "rating": 6, "units": [{"i": 0, "phase": "init", "chg": null, "v": "positive", "o": "clarity", "t": "Positive assessment of the paper's comprehensive engagement with its components and consistent empirical gains."}, {"i": 1, "phase": "init", "chg": null, "v": "negative", "o": "empirical_scope", "t": "Negative assessment of the architectural approach as potentially limited."}, {"i": 2, "phase": "init", "chg": null, "v": "negative", "o": "empirical_scope", "t": "Negative assessment of dataset coverage limitations."}, {"i": 3, "phase": "init", "chg": null, "v": "negative", "o": "robustness_sensitivity", "t": "Negative assessment due to unexplained trade-offs."}, {"i": 4, "phase": "init", "chg": null, "v": "negative", "o": "empirical_scope", "t": "Negative assessment due to incomplete contextualization."}, {"i": 5, "phase": "init", "chg": null, "v": "negative", "o": "stats_metrics", "t": "Negative assessment regarding statistical completeness."}, {"i": 6, "phase": "init", "chg": null, "v": "negative", "o": "method_design", "t": "Questioning the specificity of the learned features."}, {"i": 7, "phase": "init", "chg": null, "v": "negative", "o": "empirical_scope", "t": "Questioning the breadth of generalization."}, {"i": 8, "phase": "init", "chg": null, "v": "negative", "o": "method_design", "t": "Questioning architectural consistency."}, {"i": 9, "phase": "init", "chg": null, "v": "positive", "o": "empirical_scope", "t": "Positive interest in interpretability aspects."}, {"i": 10, "phase": "post", "chg": "strengthened", "v": "positive", "o": "baselines_ablations", "t": "Updated positive assessment regarding dataset coverage robustness."}, {"i": 11, "phase": "post", "chg": "strengthened", "v": "positive", "o": "theory", "t": "Updated positive assessment regarding generalization capabilities."}, {"i": 12, "phase": "post", "chg": "strengthened", "v": "positive", "o": "theory", "t": "Updated assessment that provides a plausible alternative explanation for trade-offs, though not a complete mechanistic account."}, {"i": 13, "phase": "post", "chg": "strengthened", "v": "positive", "o": "stats_metrics", "t": "Updated positive assessment regarding baseline comparisons."}, {"i": 14, "phase": "post", "chg": "strengthened", "v": "positive", "o": "stats_metrics", "t": "Updated positive assessment regarding statistical reporting."}, {"i": 15, "phase": "post", "chg": "strengthened", "v": "positive", "o": "empirical_scope", "t": "Updated positive assessment regarding architectural consistency."}, {"i": 16, "phase": "post", "chg": "strengthened", "v": "positive", "o": "empirical_scope", "t": "Updated positive assessment regarding interpretability."}, {"i": 17, "phase": "post", "chg": null, "v": "negative", "o": "clarity", "t": "Neutral assessment; concern remains partially unresolved."}, {"i": 18, "phase": "post", "chg": null, "v": "mixed", "o": "stats_metrics", "t": "Neutral assessment; comparison was ultimately addressed but too late for direct influence."}]}, {"key": "DsWe", "rating": 6, "units": [{"i": 0, "phase": "init", "chg": null, "v": "positive", "o": "problem_framing", "t": "Positive assessment of the biological framing and wide-ranging experiments."}, {"i": 1, "phase": "init", "chg": null, "v": "negative", "o": "problem_framing", "t": "Negative assessment due to lack of theoretical motivation."}, {"i": 2, "phase": "init", "chg": null, "v": "negative", "o": "method_design", "t": "Negative assessment due to oversimplification of biological complexity."}, {"i": 3, "phase": "init", "chg": null, "v": "negative", "o": "empirical_scope", "t": "Questioning the depth of mechanistic evidence."}, {"i": 4, "phase": "init", "chg": null, "v": "negative", "o": "empirical_scope", "t": "Questioning generalization capacity."}, {"i": 5, "phase": "init", "chg": null, "v": "negative", "o": "empirical_scope", "t": "Questioning biological specificity."}, {"i": 6, "phase": "post", "chg": "strengthened", "v": "mixed", "o": "problem_framing", "t": "Updated assessment that provides some mechanistic insight but acknowledges remaining uncertainty."}, {"i": 7, "phase": "post", "chg": "strengthened", "v": "positive", "o": "clarity", "t": "Updated positive assessment regarding labeling scheme clarity."}, {"i": 8, "phase": "post", "chg": "strengthened", "v": "positive", "o": "theory", "t": "Updated positive assessment regarding generalization."}]}, {"key": "MQMA", "rating": 6, "units": [{"i": 0, "phase": "init", "chg": null, "v": "positive", "o": "compute_cost", "t": "Positive assessment of the dataset's value and breadth."}, {"i": 1, "phase": "init", "chg": null, "v": "negative", "o": "baselines_ablations", "t": "Negative assessment of dataset coverage limitations."}, {"i": 2, "phase": "init", "chg": null, "v": "negative", "o": "compute_cost", "t": "Negative assessment of methodological exploration gaps."}, {"i": 3, "phase": "post", "chg": "strengthened", "v": "positive", "o": "baselines_ablations", "t": "Updated positive assessment regarding dataset coverage."}, {"i": 4, "phase": "post", "chg": "strengthened", "v": "positive", "o": "compute_cost", "t": "Updated positive assessment regarding methodological thoroughness."}]}, {"key": "RdsB", "rating": 4, "units": [{"i": 0, "phase": "init", "chg": null, "v": "negative", "o": "empirical_scope", "t": "Negative assessment due to insufficient dataset description."}, {"i": 1, "phase": "init", "chg": null, "v": "negative", "o": "empirical_scope", "t": "Strong negative assessment regarding experimental validity."}, {"i": 2, "phase": "init", "chg": null, "v": "negative", "o": "baselines_ablations", "t": "Negative assessment due to unexplained methodological choice."}, {"i": 3, "phase": "init", "chg": null, "v": "negative", "o": "empirical_scope", "t": "Questioning experimental setup validity."}, {"i": 4, "phase": "post", "chg": "strengthened", "v": "positive", "o": "empirical_scope", "t": "Updated positive assessment regarding dataset description."}, {"i": 5, "phase": "post", "chg": "strengthened", "v": "positive", "o": "empirical_scope", "t": "Updated positive assessment regarding experimental validity."}, {"i": 6, "phase": "post", "chg": "strengthened", "v": "positive", "o": "baselines_ablations", "t": "Updated positive assessment regarding methodological transparency."}, {"i": 7, "phase": "post", "chg": null, "v": "mixed", "o": "empirical_scope", "t": "Conditional positive assessment pending further evidence."}, {"i": 8, "phase": "post", "chg": "strengthened", "v": "positive", "o": "robustness_sensitivity", "t": "Updated positive assessment regarding data integrity."}, {"i": 9, "phase": "post", "chg": "strengthened", "v": "positive", "o": "empirical_scope", "t": "Updated positive assessment regarding dataset characterization."}, {"i": 10, "phase": "post", "chg": "strengthened", "v": "positive", "o": "baselines_ablations", "t": "Updated positive assessment regarding manuscript updates."}]}]}, {"scenario": "an entrenchment", "forum": "apaLoTumdO", "title": "CE-Nav: Flow-Guided Reinforcement Refinement for Cross-Embodiment Local Navigation", "decision": "Accept (Poster)", "comments": 16, "reviewers": [{"key": "3iX8", "rating": 6, "units": [{"i": 0, "phase": "init", "chg": null, "v": "positive", "o": "compute_cost", "t": "The paper is solid with fixable issues."}, {"i": 1, "phase": "init", "chg": null, "v": "negative", "o": "method_design", "t": "The design has a gap regarding potential narrow environments."}, {"i": 2, "phase": "init", "chg": null, "v": "negative", "o": "method_design", "t": "The explanation of cross-embodiment variations is insufficient."}, {"i": 3, "phase": "init", "chg": null, "v": "negative", "o": "method_design", "t": "The scope of training conditions is limited."}, {"i": 4, "phase": "post", "chg": "strengthened", "v": "positive", "o": "method_design", "t": "The concern about narrow environments is mitigated by empirical evidence of large-robot success."}, {"i": 5, "phase": "post", "chg": "strengthened", "v": "positive", "o": "robustness_sensitivity", "t": "The clarification on failure modes provides a more nuanced understanding of platform differences."}, {"i": 6, "phase": "post", "chg": "strengthened", "v": "mixed", "o": "theory", "t": "The mechanism for handling dynamics is plausible, though not fully resolved for high-speed scenes."}]}, {"key": "Ko6r", "rating": 6, "units": [{"i": 0, "phase": "init", "chg": null, "v": "negative", "o": "novelty", "t": "The technical novelty is moderate and requires clearer differentiation from prior art."}, {"i": 1, "phase": "init", "chg": null, "v": "negative", "o": "related_work", "t": "The related work section is incomplete."}, {"i": 2, "phase": "init", "chg": null, "v": "negative", "o": "theory", "t": "The conceptual framing lacks precision."}, {"i": 3, "phase": "init", "chg": null, "v": "negative", "o": "compute_cost", "t": "The experimental evidence is fragmented and lacks key comparisons."}, {"i": 4, "phase": "post", "chg": "strengthened", "v": "positive", "o": "method_design", "t": "The distinction from prior residual RL methods is empirically validated."}, {"i": 5, "phase": "post", "chg": "strengthened", "v": "positive", "o": "robustness_sensitivity", "t": "The architectural choice is justified by its resistance to catastrophic interference."}, {"i": 6, "phase": "post", "chg": "strengthened", "v": "positive", "o": "method_design", "t": "The efficiency claims are substantiated by new benchmarking."}, {"i": 7, "phase": "post", "chg": "strengthened", "v": "positive", "o": "method_design", "t": "The method shows robustness to significant embodiment variations."}, {"i": 8, "phase": "post", "chg": null, "v": "mixed", "o": "method_design", "t": "The reward design improvement is plausible but lacks empirical isolation."}]}, {"key": "rjvc", "rating": 4, "units": [{"i": 0, "phase": "init", "chg": null, "v": "negative", "o": "empirical_scope", "t": "The reliance on planner-specific priors is a potential limitation."}, {"i": 1, "phase": "init", "chg": null, "v": "negative", "o": "method_design", "t": "The method's safety guarantees across diverse embodiments are unverified."}, {"i": 2, "phase": "init", "chg": null, "v": "negative", "o": "empirical_scope", "t": "The real-time feasibility for demanding applications is questionable."}, {"i": 3, "phase": "post", "chg": "strengthened", "v": "positive", "o": "method_design", "t": "The inductive bias concern is mitigated by evidence of planner-invariant learning."}, {"i": 4, "phase": "post", "chg": "strengthened", "v": "positive", "o": "baselines_ablations", "t": "The method meets real-time requirements for the tested platforms."}, {"i": 5, "phase": "post", "chg": "strengthened", "v": "positive", "o": "theory", "t": "The method is theoretically and empirically supported for embodiment mismatch."}, {"i": 6, "phase": "post", "chg": null, "v": "mixed", "o": "method_design", "t": "The scope clarification is logical, but empirical support for the slow system is weak."}]}, {"key": "oUen", "rating": 2, "units": [{"i": 0, "phase": "init", "chg": null, "v": "negative", "o": "baselines_ablations", "t": "The performance claims are undermined by inadequate baselines."}, {"i": 1, "phase": "init", "chg": null, "v": "negative", "o": "theory", "t": "The architectural complexity is unjustified by demonstrated benefits."}, {"i": 2, "phase": "init", "chg": null, "v": "negative", "o": "novelty", "t": "The contribution is marginal and overlaps significantly with prior art."}, {"i": 3, "phase": "init", "chg": null, "v": "negative", "o": "method_design", "t": "The method fails to meet the paradigm of zero-shot embodiment transfer."}, {"i": 4, "phase": "init", "chg": null, "v": "negative", "o": "stats_metrics", "t": "The results are weak and lack compelling context."}, {"i": 5, "phase": "post", "chg": "strengthened", "v": "positive", "o": "baselines_ablations", "t": "The method achieves state-of-the-art performance."}, {"i": 6, "phase": "post", "chg": "strengthened", "v": "positive", "o": "compute_cost", "t": "The architectural complexity is necessary to avoid catastrophic failures."}, {"i": 7, "phase": "post", "chg": "strengthened", "v": "positive", "o": "novelty", "t": "The hybrid loss offers a distinct advantage over distillation-based methods."}, {"i": 8, "phase": "post", "chg": "strengthened", "v": "positive", "o": "baselines_ablations", "t": "The need for adaptation is justified by physical realities."}, {"i": 9, "phase": "post", "chg": "strengthened", "v": "positive", "o": "method_design", "t": "The performance is stronger than initial metrics suggested."}, {"i": 10, "phase": "post", "chg": "strengthened", "v": "positive", "o": "baselines_ablations", "t": "The choice of CNF is empirically justified."}, {"i": 11, "phase": "post", "chg": "strengthened", "v": "positive", "o": "theory", "t": "The velocity-control approach is superior for this application."}]}]}, {"scenario": "a reversal", "forum": "OlidGx8oKr", "title": "Incentive-Aligned Multi-Source LLM Summaries", "decision": "Accept (Poster)", "comments": 19, "reviewers": [{"key": "Reviewer_fpEF", "rating": 6, "units": [{"i": 0, "phase": "init", "chg": null, "v": "positive", "o": "method_design", "t": "The contribution is valuable and addresses a relevant problem."}, {"i": 1, "phase": "init", "chg": null, "v": "positive", "o": "stats_metrics", "t": "The theoretical framework is strong."}, {"i": 2, "phase": "init", "chg": null, "v": "positive", "o": "method_design", "t": "The experimental results are convincing."}, {"i": 3, "phase": "init", "chg": null, "v": "negative", "o": "empirical_scope", "t": "Evaluation is limited by small scale and artificial data."}, {"i": 4, "phase": "init", "chg": null, "v": "negative", "o": "compute_cost", "t": "Practicality is hindered by high computational cost."}, {"i": 5, "phase": "init", "chg": null, "v": "negative", "o": "clarity", "t": "Presentation is poor and needs simplification."}, {"i": 6, "phase": "init", "chg": null, "v": "negative", "o": "empirical_scope", "t": "Potential trade-off between incentive alignment and information diversity exists."}, {"i": 7, "phase": "init", "chg": null, "v": "negative", "o": "theory", "t": "Theoretical assumption of exchangeability is practically flawed."}, {"i": 8, "phase": "post", "chg": "reversed", "v": "positive", "o": "compute_cost", "t": "Cost concern is addressed."}, {"i": 9, "phase": "post", "chg": "reversed", "v": "positive", "o": "stats_metrics", "t": "Scale concern is mitigated."}, {"i": 10, "phase": "post", "chg": "reversed", "v": "positive", "o": "empirical_scope", "t": "Validation is strengthened by real-world data."}, {"i": 11, "phase": "post", "chg": "reversed", "v": "positive", "o": "stats_metrics", "t": "Presentation and theoretical flexibility improved."}, {"i": 12, "phase": "post", "chg": "reversed", "v": "positive", "o": "baselines_ablations", "t": "Diversity concern is addressed."}]}, {"key": "Reviewer_eWFW", "rating": 4, "units": [{"i": 0, "phase": "init", "chg": null, "v": "positive", "o": "problem_framing", "t": "The core idea is sound."}, {"i": 1, "phase": "init", "chg": null, "v": "positive", "o": "empirical_scope", "t": "The experimental results support the claims."}, {"i": 2, "phase": "init", "chg": null, "v": "negative", "o": "stats_metrics", "t": "Empirical validation is insufficient for the claims made regarding robustness."}, {"i": 3, "phase": "init", "chg": null, "v": "negative", "o": "stats_metrics", "t": "The evaluation methodology is incomplete for assessing factual reliability."}]}, {"key": "Reviewer_hbFR", "rating": 4, "units": [{"i": 0, "phase": "init", "chg": null, "v": "positive", "o": "method_design", "t": "The conceptual approach is strong and theoretically grounded."}, {"i": 1, "phase": "init", "chg": null, "v": "negative", "o": "theory", "t": "Fundamental problems with the theoretical assumptions reduce soundness."}, {"i": 2, "phase": "init", "chg": null, "v": "negative", "o": "compute_cost", "t": "Practical feasibility is questionable due to high computational cost."}, {"i": 3, "phase": "init", "chg": null, "v": "negative", "o": "baselines_ablations", "t": "Evaluation is incomplete due to lack of strong external comparators."}, {"i": 4, "phase": "post", "chg": "reversed", "v": "positive", "o": "compute_cost", "t": "The computational concern is substantially mitigated."}, {"i": 5, "phase": "post", "chg": "reversed", "v": "positive", "o": "theory", "t": "Theoretical concerns are contextualized and less critical."}, {"i": 6, "phase": "post", "chg": null, "v": "mixed", "o": "baselines_ablations", "t": "Concerns about baselines are clarified but not fully resolved."}, {"i": 7, "phase": "post", "chg": "reversed", "v": "positive", "o": "empirical_scope", "t": "Empirical validation is much stronger."}]}]}, {"scenario": "a reversal", "forum": "5fCDEz43ya", "title": "Token-Guard: Towards Token-Level Hallucination Control via Self-Checking Decoding", "decision": "Accept (Poster)", "comments": 16, "reviewers": [{"key": "35wE", "rating": 8, "units": [{"i": 0, "phase": "init", "chg": null, "v": "positive", "o": "method_design", "t": "Positive assessment of conceptual contribution."}, {"i": 1, "phase": "init", "chg": null, "v": "positive", "o": "compute_cost", "t": "Strong positive assessment of results."}, {"i": 2, "phase": "init", "chg": null, "v": "positive", "o": "empirical_scope", "t": "Positive assessment of methodological coherence."}, {"i": 3, "phase": "init", "chg": null, "v": "negative", "o": "compute_cost", "t": "Concern about misleading framing."}, {"i": 4, "phase": "init", "chg": null, "v": "negative", "o": "robustness_sensitivity", "t": "Concern about reproducibility and robustness."}, {"i": 5, "phase": "init", "chg": null, "v": "negative", "o": "stats_metrics", "t": "Epistemic concern about foundational assumptions."}, {"i": 6, "phase": "init", "chg": null, "v": "negative", "o": "empirical_scope", "t": "Concern about limited evaluation scope."}, {"i": 7, "phase": "post", "chg": null, "v": "positive", "o": "related_work", "t": "Maintained positive rating after clarification."}]}, {"key": "5cxV", "rating": 6, "units": [{"i": 0, "phase": "init", "chg": null, "v": "positive", "o": "method_design", "t": "Positive evaluation of the core mechanism's utility."}, {"i": 1, "phase": "init", "chg": null, "v": "positive", "o": "compute_cost", "t": "Positive assessment of practical accessibility."}, {"i": 2, "phase": "init", "chg": null, "v": "positive", "o": "baselines_ablations", "t": "Positive assessment of empirical validity."}, {"i": 3, "phase": "init", "chg": null, "v": "negative", "o": "stats_metrics", "t": "Concern regarding practical applicability."}, {"i": 4, "phase": "init", "chg": null, "v": "negative", "o": "baselines_ablations", "t": "Skepticism about unsubstantiated comparative claims."}, {"i": 5, "phase": "init", "chg": null, "v": "negative", "o": "empirical_scope", "t": "Concern about boundary conditions and scope limitations."}, {"i": 6, "phase": "init", "chg": null, "v": "negative", "o": "novelty", "t": "Factual challenge to the paper's originality framing."}]}, {"key": "ATsM", "rating": 4, "units": [{"i": 0, "phase": "init", "chg": null, "v": "positive", "o": "clarity", "t": "Positive assessment of presentation quality."}, {"i": 1, "phase": "init", "chg": null, "v": "positive", "o": "baselines_ablations", "t": "Positive assessment of results."}, {"i": 2, "phase": "init", "chg": null, "v": "negative", "o": "robustness_sensitivity", "t": "Concern about practical utility and potential overfitting."}, {"i": 3, "phase": "init", "chg": null, "v": "negative", "o": "empirical_scope", "t": "Technical concern about method logic."}, {"i": 4, "phase": "init", "chg": null, "v": "negative", "o": "compute_cost", "t": "Concern about resource efficiency."}, {"i": 5, "phase": "init", "chg": null, "v": "negative", "o": "stats_metrics", "t": "Critical concern about experimental integrity."}, {"i": 6, "phase": "init", "chg": null, "v": "negative", "o": "method_design", "t": "Concern about theoretical depth."}, {"i": 7, "phase": "init", "chg": null, "v": "negative", "o": "theory", "t": "Concern about algorithmic necessity."}, {"i": 8, "phase": "init", "chg": null, "v": "negative", "o": "baselines_ablations", "t": "Concern about incomplete evaluation."}]}, {"key": "eh4g", "rating": 4, "units": [{"i": 0, "phase": "init", "chg": null, "v": "positive", "o": "empirical_scope", "t": "Neutral-to-positive acknowledgment of structure."}, {"i": 1, "phase": "init", "chg": null, "v": "positive", "o": "empirical_scope", "t": "Positive assessment of effort."}, {"i": 2, "phase": "init", "chg": "reversed", "v": "negative", "o": "novelty", "t": "Strong negative assessment of novelty and completeness."}, {"i": 3, "phase": "init", "chg": "reversed", "v": "negative", "o": "reproducibility", "t": "Negative assessment of transparency."}, {"i": 4, "phase": "init", "chg": "reversed", "v": "negative", "o": "stats_metrics", "t": "Strong negative assessment of evaluation validity."}, {"i": 5, "phase": "init", "chg": "reversed", "v": "negative", "o": "empirical_scope", "t": "Concern about generalizability."}, {"i": 6, "phase": "init", "chg": "reversed", "v": "negative", "o": "robustness_sensitivity", "t": "Concern about incomplete risk assessment."}, {"i": 7, "phase": "post", "chg": "reversed", "v": "positive", "o": "novelty", "t": "Raised score due to resolved concerns."}]}]}, {"scenario": "a split verdict", "forum": "lR8GufFQMb", "title": "Test-time scaling of diffusions with flow maps", "decision": "Reject", "comments": 10, "reviewers": [{"key": "sdVa", "rating": 10, "units": [{"i": 0, "phase": "init", "chg": null, "v": "positive", "o": "theory", "t": "Solid contribution with no major weaknesses."}, {"i": 1, "phase": "init", "chg": null, "v": "positive", "o": "baselines_ablations", "t": "Excellent experimental strength."}, {"i": 2, "phase": "init", "chg": null, "v": "negative", "o": "clarity", "t": "Minor presentational issue."}, {"i": 3, "phase": "init", "chg": null, "v": "negative", "o": "clarity", "t": "Minor presentational issue."}, {"i": 4, "phase": "init", "chg": null, "v": "negative", "o": "clarity", "t": "Presentational improvement needed for clarity."}, {"i": 5, "phase": "post", "chg": null, "v": "negative", "o": "clarity", "t": "Requested evidence missing/mislocated."}]}, {"key": "84bA", "rating": 6, "units": [{"i": 0, "phase": "init", "chg": null, "v": "positive", "o": "empirical_scope", "t": "Highly original, significant, and well-executed contribution."}, {"i": 1, "phase": "init", "chg": null, "v": "negative", "o": "empirical_scope", "t": "Weakness due to limited applicability and cost."}, {"i": 2, "phase": "init", "chg": null, "v": "negative", "o": "method_design", "t": "Uncertain due to missing evidence."}, {"i": 3, "phase": "init", "chg": null, "v": "negative", "o": "empirical_scope", "t": "Questioned robustness."}, {"i": 4, "phase": "init", "chg": null, "v": "negative", "o": "method_design", "t": "Unclear compatibility."}, {"i": 5, "phase": "init", "chg": null, "v": "negative", "o": "empirical_scope", "t": "Uncertain generalization."}]}, {"key": "t4do", "rating": 4, "units": [{"i": 0, "phase": "init", "chg": null, "v": "positive", "o": "problem_framing", "t": "Important problem with non-trivial, theoretically principled technique."}, {"i": 1, "phase": "init", "chg": null, "v": "negative", "o": "theory", "t": "Unclear motivation."}, {"i": 2, "phase": "init", "chg": null, "v": "negative", "o": "clarity", "t": "Poor exposition hinders understanding."}, {"i": 3, "phase": "init", "chg": null, "v": "negative", "o": "stats_metrics", "t": "Insufficient evaluations."}, {"i": 4, "phase": "init", "chg": null, "v": "negative", "o": "robustness_sensitivity", "t": "Missing critical scaling evidence."}, {"i": 5, "phase": "post", "chg": null, "v": "positive", "o": "empirical_scope", "t": "Open to increasing score if new evidence is integrated."}]}, {"key": "hwMk", "rating": 0, "units": [{"i": 0, "phase": "init", "chg": null, "v": "positive", "o": "theory", "t": "Theoretical derivations are correct."}, {"i": 1, "phase": "init", "chg": null, "v": "negative", "o": "novelty", "t": "Lacks novelty; over-marketed as fundamentally new."}, {"i": 2, "phase": "init", "chg": null, "v": "negative", "o": "stats_metrics", "t": "Theoretical section is superficial regarding critical issues."}, {"i": 3, "phase": "init", "chg": null, "v": "negative", "o": "problem_framing", "t": "Presentation is misleading and lacks sobriety."}, {"i": 4, "phase": "init", "chg": null, "v": "negative", "o": "problem_framing", "t": "Weak motivation."}, {"i": 5, "phase": "init", "chg": null, "v": "negative", "o": "baselines_ablations", "t": "Empirical evidence is unconvincing and cherry-picked."}, {"i": 6, "phase": "init", "chg": null, "v": "negative", "o": "empirical_scope", "t": "Efficiency claims unsupported."}, {"i": 7, "phase": "init", "chg": null, "v": "negative", "o": "theory", "t": "Framing is intellectually dishonest."}]}]}, {"scenario": "a split verdict", "forum": "072P11r1wu", "title": "Understanding Generalization in Transformers: Error Bounds and Training Dynamics Under Benign an", "decision": "Reject", "comments": 0, "reviewers": [{"key": "ADv3", "rating": 10, "units": [{"i": 0, "phase": "init", "chg": null, "v": "positive", "o": "method_design", "t": "The fine-grained analysis is a strength."}, {"i": 1, "phase": "init", "chg": null, "v": "positive", "o": "empirical_scope", "t": "The handling of label noise is a strength."}, {"i": 2, "phase": "init", "chg": null, "v": "positive", "o": "method_design", "t": "The stage distinction is interesting and insightful."}, {"i": 3, "phase": "init", "chg": null, "v": "positive", "o": "stats_metrics", "t": "The experiments are a strength."}, {"i": 4, "phase": "init", "chg": null, "v": "positive", "o": "theory", "t": "No major concern, but weakening orthogonality is a valuable future direction."}, {"i": 5, "phase": "init", "chg": null, "v": "positive", "o": "theory", "t": "No major concern, but extending sequence length is a relevant question."}, {"i": 6, "phase": "init", "chg": null, "v": "positive", "o": "method_design", "t": "The practical scope of the theory in LMs is an open question."}]}, {"key": "VQ35", "rating": 2, "units": [{"i": 0, "phase": "init", "chg": null, "v": "negative", "o": "stats_metrics", "t": "The experimental setup is too small and synthetic to support the paper's claims."}, {"i": 1, "phase": "init", "chg": null, "v": "negative", "o": "clarity", "t": "Presentation is unclear."}, {"i": 2, "phase": "init", "chg": null, "v": "negative", "o": "theory", "t": "Contribution 1 is not a valid contribution claim."}, {"i": 3, "phase": "init", "chg": null, "v": "negative", "o": "theory", "t": "Presentation is incomplete."}, {"i": 4, "phase": "init", "chg": null, "v": "negative", "o": "empirical_scope", "t": "The claim is unexplained and potentially inaccurate."}, {"i": 5, "phase": "init", "chg": null, "v": "negative", "o": "theory", "t": "The theoretical bounds are quantitatively inaccurate."}, {"i": 6, "phase": "init", "chg": null, "v": "negative", "o": "related_work", "t": "Literature coverage is insufficient."}]}, {"key": "NCWG", "rating": 2, "units": [{"i": 0, "phase": "init", "chg": null, "v": "negative", "o": "theory", "t": "Theorem 14 is nearly vacuous."}, {"i": 1, "phase": "init", "chg": null, "v": "negative", "o": "theory", "t": "The proof of Theorem 20 is incomplete due to a threshold mismatch."}, {"i": 2, "phase": "init", "chg": null, "v": "negative", "o": "theory", "t": "The analysis ignores necessary correlation terms, rendering the bound invalid."}, {"i": 3, "phase": "init", "chg": null, "v": "negative", "o": "theory", "t": "The transfer of equations from clean-label to noisy-label settings is unjustified and likely incorrect."}, {"i": 4, "phase": "init", "chg": null, "v": "negative", "o": "theory", "t": "The Gaussian assumption is questionable beyond $t=0$."}, {"i": 5, "phase": "init", "chg": null, "v": "negative", "o": "theory", "t": "The novelty of the stage decomposition is unclear."}]}, {"key": "3Ed7", "rating": 0, "units": [{"i": 0, "phase": "init", "chg": null, "v": "negative", "o": "theory", "t": "The proofs are incorrect/incomplete because the transfer of equations from a clean-label setting to a noisy-label setting is unjustified."}, {"i": 1, "phase": "init", "chg": null, "v": "negative", "o": "clarity", "t": "The claim lacks formal precision."}, {"i": 2, "phase": "init", "chg": null, "v": "negative", "o": "theory", "t": "The validity of the stage-one analysis beyond initialization is unverified."}, {"i": 3, "phase": "init", "chg": null, "v": "negative", "o": "theory", "t": "The normality claim is statistically invalid due to dependence between weights and noise."}]}]}, {"scenario": "a split verdict", "forum": "zvw9Hiwa0i", "title": "Beyond Scattered Acceptance: Fast and Coherent Inference for DLMs via Longest Stable Prefixes", "decision": "Accept (Poster)", "comments": 8, "reviewers": [{"key": "Cvwn", "rating": 10, "units": [{"i": 0, "phase": "init", "chg": null, "v": "positive", "o": "stats_metrics", "t": "The design is concrete and well-engineered."}, {"i": 1, "phase": "init", "chg": null, "v": "positive", "o": "compute_cost", "t": "The topology is elegant."}, {"i": 2, "phase": "init", "chg": null, "v": "positive", "o": "baselines_ablations", "t": "The empirical support is adequate."}, {"i": 3, "phase": "init", "chg": null, "v": "negative", "o": "theory", "t": "This is a weakness regarding robustness."}, {"i": 4, "phase": "init", "chg": null, "v": "negative", "o": "problem_framing", "t": "The theoretical framing is light."}, {"i": 5, "phase": "init", "chg": null, "v": "negative", "o": "compute_cost", "t": "The evaluation is incomplete regarding error correction overhead."}, {"i": 6, "phase": "post", "chg": null, "v": "negative", "o": "theory", "t": "The brittleness remains an unresolved limitation."}, {"i": 7, "phase": "post", "chg": null, "v": "mixed", "o": "compute_cost", "t": "The argument reasonably addresses the concern about repair costs, though it measures active suffix stability rather than committed token repair directly."}, {"i": 8, "phase": "post", "chg": null, "v": "negative", "o": "problem_framing", "t": "The theoretical gap remains unfilled."}]}, {"key": "Yt8W", "rating": 4, "units": [{"i": 0, "phase": "init", "chg": null, "v": "negative", "o": "method_design", "t": "The novelty is thin/incremental."}, {"i": 1, "phase": "init", "chg": null, "v": "negative", "o": "theory", "t": "The method suffers from limited flexibility."}, {"i": 2, "phase": "init", "chg": null, "v": "negative", "o": "theory", "t": "Key parameters lack grounding."}, {"i": 3, "phase": "init", "chg": null, "v": "negative", "o": "compute_cost", "t": "The evaluation benchmarks are insufficiently challenging."}, {"i": 4, "phase": "init", "chg": null, "v": "negative", "o": "compute_cost", "t": "The speedup analysis is opaque."}, {"i": 5, "phase": "init", "chg": null, "v": "negative", "o": "empirical_scope", "t": "The evaluation lacks robustness checks."}, {"i": 6, "phase": "post", "chg": null, "v": "mixed", "o": "novelty", "t": "The novelty concern is reframed but not fully resolved."}, {"i": 7, "phase": "post", "chg": null, "v": "positive", "o": "related_work", "t": "The flexibility limitation is acknowledged and documented."}, {"i": 8, "phase": "post", "chg": null, "v": "positive", "o": "theory", "t": "The lack of grounding is addressed by design clarification."}, {"i": 9, "phase": "post", "chg": null, "v": "mixed", "o": "stats_metrics", "t": "The benchmark concern is partially addressed but highlights evaluation limitations."}, {"i": 10, "phase": "post", "chg": null, "v": "positive", "o": "compute_cost", "t": "The speedup attribution is clarified."}, {"i": 11, "phase": "post", "chg": null, "v": "positive", "o": "empirical_scope", "t": "The robustness concern is addressed."}]}, {"key": "kHHK", "rating": 4, "units": [{"i": 0, "phase": "init", "chg": null, "v": "negative", "o": "baselines_ablations", "t": "The evaluation lacks necessary comparative baselines."}, {"i": 1, "phase": "init", "chg": null, "v": "negative", "o": "robustness_sensitivity", "t": "The method's hyperparameter stability is unverified."}, {"i": 2, "phase": "init", "chg": null, "v": "negative", "o": "clarity", "t": "The figure contains a critical error inconsistent with the algorithm."}, {"i": 3, "phase": "post", "chg": null, "v": "mixed", "o": "baselines_ablations", "t": "The baseline concern is only partially addressed."}, {"i": 4, "phase": "post", "chg": null, "v": "positive", "o": "robustness_sensitivity", "t": "The hyperparameter sensitivity concern is convincingly addressed."}, {"i": 5, "phase": "post", "chg": null, "v": "positive", "o": "clarity", "t": "The figure error is corrected."}]}, {"key": "jKnz", "rating": 2, "units": [{"i": 0, "phase": "init", "chg": null, "v": "negative", "o": "novelty", "t": "The contribution is fundamentally misframed."}, {"i": 1, "phase": "init", "chg": null, "v": "negative", "o": "theory", "t": "The proposed enhancements are essential fixes, not novel additions."}, {"i": 2, "phase": "init", "chg": null, "v": "negative", "o": "empirical_scope", "t": "The evaluation lacks necessary diversity."}, {"i": 3, "phase": "init", "chg": null, "v": "negative", "o": "theory", "t": "The method loses a key DLM advantage."}, {"i": 4, "phase": "init", "chg": null, "v": "negative", "o": "compute_cost", "t": "There is a technical gap in the caching explanation."}, {"i": 5, "phase": "init", "chg": null, "v": "negative", "o": "method_design", "t": "The evaluation is incomplete regarding complex generation scenarios."}, {"i": 6, "phase": "post", "chg": null, "v": "negative", "o": "theory", "t": "The categorical concern remains unresolved; the response is a narrowing of scope rather than a refutation."}, {"i": 7, "phase": "post", "chg": null, "v": "positive", "o": "compute_cost", "t": "The concern is addressed empirically."}, {"i": 8, "phase": "post", "chg": null, "v": "negative", "o": "problem_framing", "t": "The trade-off is acknowledged but the fundamental category error remains."}, {"i": 9, "phase": "post", "chg": null, "v": "mixed", "o": "reproducibility", "t": "The technical gap is mitigated by citation but not fully resolved by original analysis."}, {"i": 10, "phase": "post", "chg": null, "v": "mixed", "o": "problem_framing", "t": "The concern is partially addressed, but the validation method is weak."}]}]}, {"scenario": "unanimity", "forum": "5LMdnUdAoy", "title": "Difficult Examples Hurt Unsupervised Contrastive Learning: A Theoretical Perspective", "decision": "Accept (Oral)", "comments": 8, "reviewers": [{"key": "Vvxq", "rating": 6, "units": [{"i": 0, "phase": "init", "chg": null, "v": "negative", "o": "stats_metrics", "t": "Experiments are only partially sound due to insufficient statistical power."}, {"i": 1, "phase": "init", "chg": null, "v": "negative", "o": "empirical_scope", "t": "Potential confound identified that weakens causal claim for TinyImagenet results."}, {"i": 2, "phase": "init", "chg": null, "v": "negative", "o": "related_work", "t": "Prior work engagement is incomplete/superficial."}, {"i": 3, "phase": "init", "chg": null, "v": "negative", "o": "empirical_scope", "t": "Theoretical claims face tension with existing empirical evidence regarding hard negatives."}, {"i": 4, "phase": "init", "chg": null, "v": "negative", "o": "baselines_ablations", "t": "Methodology description is insufficiently detailed."}, {"i": 5, "phase": "init", "chg": null, "v": "negative", "o": "theory", "t": "Example consistency with theory is ambiguous."}, {"i": 6, "phase": "init", "chg": null, "v": "negative", "o": "empirical_scope", "t": "Theoretical model granularity is unspecified."}, {"i": 7, "phase": "init", "chg": null, "v": "negative", "o": "theory", "t": "Notation contains errors or confusing definitions."}, {"i": 8, "phase": "init", "chg": null, "v": "negative", "o": "theory", "t": "Lack of visualization reduces accessibility and verification ease."}, {"i": 9, "phase": "init", "chg": null, "v": "negative", "o": "empirical_scope", "t": "Validation is limited to real datasets; theoretical purity is not tested."}, {"i": 10, "phase": "post", "chg": null, "v": "positive", "o": "robustness_sensitivity", "t": "Statistical concern partially mitigated by new evidence of low variance."}, {"i": 11, "phase": "post", "chg": null, "v": "mixed", "o": "theory", "t": "Confound explanation is plausible but shifts the narrative rather than eliminating the variable."}, {"i": 12, "phase": "post", "chg": null, "v": "positive", "o": "theory", "t": "Distinction clarifies the theoretical landscape."}, {"i": 13, "phase": "post", "chg": null, "v": "positive", "o": "robustness_sensitivity", "t": "Clarity improved."}, {"i": 14, "phase": "post", "chg": null, "v": "positive", "o": "clarity", "t": "Notation confusion resolved."}, {"i": 15, "phase": "post", "chg": null, "v": "negative", "o": "empirical_scope", "t": "Request rejected with justification; synthetic validation remains absent."}, {"i": 16, "phase": "post", "chg": null, "v": "positive", "o": "clarity", "t": "Presentation issue addressed."}, {"i": 17, "phase": "post", "chg": null, "v": "mixed", "o": "related_work", "t": "Justification provided but does not fully address the nuance gap."}, {"i": 18, "phase": "post", "chg": null, "v": "positive", "o": "empirical_scope", "t": "Visualization request fulfilled."}, {"i": 19, "phase": "post", "chg": null, "v": "mixed", "o": "stats_metrics", "t": "Partial clarification; granularity limitation remains."}, {"i": 20, "phase": "post", "chg": null, "v": "positive", "o": "stats_metrics", "t": "Initial judgment stands."}]}, {"key": "Dph8", "rating": 6, "units": [{"i": 0, "phase": "init", "chg": null, "v": "negative", "o": "theory", "t": "Theoretical assumptions are potentially incompatible with cosine similarity application."}, {"i": 1, "phase": "init", "chg": null, "v": "negative", "o": "theory", "t": "Bounds are of limited value without tightness context."}, {"i": 2, "phase": "init", "chg": null, "v": "negative", "o": "empirical_scope", "t": "Empirical evidence is weak due to modest effects and low run count."}, {"i": 3, "phase": "init", "chg": null, "v": "negative", "o": "theory", "t": "Scope of contribution is unclear beyond vision."}, {"i": 4, "phase": "init", "chg": null, "v": "negative", "o": "clarity", "t": "Presentation quality is poor."}, {"i": 5, "phase": "post", "chg": null, "v": "positive", "o": "empirical_scope", "t": "Theoretical assumption concern resolved."}, {"i": 6, "phase": "post", "chg": null, "v": "positive", "o": "theory", "t": "Bound tightness concern partially addressed."}, {"i": 7, "phase": "post", "chg": null, "v": "positive", "o": "method_design", "t": "Scope concern addressed."}, {"i": 8, "phase": "post", "chg": null, "v": "positive", "o": "clarity", "t": "Presentation issues addressed."}, {"i": 9, "phase": "post", "chg": null, "v": "positive", "o": "stats_metrics", "t": "Initial judgment stands."}]}, {"key": "3ZD8", "rating": 6, "units": [{"i": 0, "phase": "init", "chg": null, "v": "positive", "o": "novelty", "t": "Finding is novel and empirically supported."}, {"i": 1, "phase": "init", "chg": null, "v": "positive", "o": "theory", "t": "Theoretical framework is strong."}, {"i": 2, "phase": "init", "chg": null, "v": "positive", "o": "problem_framing", "t": "Solutions are well-grounded and effective."}, {"i": 3, "phase": "init", "chg": null, "v": "positive", "o": "empirical_scope", "t": "Validation is comprehensive."}, {"i": 4, "phase": "init", "chg": null, "v": "negative", "o": "clarity", "t": "Core concept definition is circular and potentially flawed."}, {"i": 5, "phase": "init", "chg": null, "v": "negative", "o": "compute_cost", "t": "Scalability is questionable."}, {"i": 6, "phase": "init", "chg": null, "v": "negative", "o": "robustness_sensitivity", "t": "Theory-practice connection is weak."}, {"i": 7, "phase": "init", "chg": null, "v": "negative", "o": "theory", "t": "Theoretical conditions may be too restrictive for practice."}, {"i": 8, "phase": "init", "chg": null, "v": "negative", "o": "theory", "t": "Claimed distinction lacks empirical support."}, {"i": 9, "phase": "post", "chg": null, "v": "mixed", "o": "stats_metrics", "t": "Circularity concern partially addressed but not fully resolved."}, {"i": 10, "phase": "post", "chg": null, "v": "mixed", "o": "compute_cost", "t": "Scalability concern partially addressed."}, {"i": 11, "phase": "post", "chg": null, "v": "mixed", "o": "robustness_sensitivity", "t": "Theory-practice gap acknowledged but not closed."}, {"i": 12, "phase": "post", "chg": null, "v": "positive", "o": "theory", "t": "Theorem condition concern resolved."}, {"i": 13, "phase": "post", "chg": null, "v": "positive", "o": "stats_metrics", "t": "Distinction concern addressed via empirical evidence."}, {"i": 14, "phase": "post", "chg": null, "v": "positive", "o": "stats_metrics", "t": "Initial judgment stands."}]}]}, {"scenario": "unanimity", "forum": "XhqoDBouWS", "title": "Why Attention Patterns Exist: A Unifying Temporal Perspective Analysis", "decision": "Accept (Poster)", "comments": 24, "reviewers": [{"key": "Qj6A", "rating": 4, "units": [{"i": 0, "phase": "init", "chg": null, "v": "negative", "o": "theory", "t": "The theoretical foundation requires explicit scope definition to be valid."}, {"i": 1, "phase": "init", "chg": null, "v": "negative", "o": "empirical_scope", "t": "The methodological robustness regarding the similarity metric is unverified."}, {"i": 2, "phase": "init", "chg": null, "v": "negative", "o": "theory", "t": "The theoretical claims lack sufficient mathematical validity."}, {"i": 3, "phase": "init", "chg": null, "v": "negative", "o": "empirical_scope", "t": "The pruning strategy's safety regarding critical stable patterns is questionable."}, {"i": 4, "phase": "post", "chg": null, "v": "positive", "o": "theory", "t": "The scope clarification adequately resolves the initial worry about universality."}, {"i": 5, "phase": "post", "chg": null, "v": "positive", "o": "stats_metrics", "t": "The metric choice is empirically justified."}, {"i": 6, "phase": "post", "chg": null, "v": "positive", "o": "clarity", "t": "The mathematical rigor has been restored through rigorous correction."}, {"i": 7, "phase": "post", "chg": null, "v": "positive", "o": "method_design", "t": "The pruning logic is safe for critical patterns."}, {"i": 8, "phase": "post", "chg": null, "v": "positive", "o": "clarity", "t": "Positive evaluation update."}, {"i": 9, "phase": "post", "chg": null, "v": "negative", "o": "theory", "t": "A residual mathematical error remains in the correction."}, {"i": 10, "phase": "post", "chg": null, "v": "positive", "o": "theory", "t": "The residual error is resolved."}]}, {"key": "V8Hz", "rating": 4, "units": [{"i": 0, "phase": "init", "chg": null, "v": "negative", "o": "method_design", "t": "The conceptual contribution is overstated and lacks clarity regarding the joint effect."}, {"i": 1, "phase": "init", "chg": null, "v": "negative", "o": "theory", "t": "The theoretical validity is weakened by unsupported assumptions."}, {"i": 2, "phase": "init", "chg": null, "v": "negative", "o": "baselines_ablations", "t": "The empirical evaluation is incomplete relative to state-of-the-art."}, {"i": 3, "phase": "init", "chg": null, "v": "negative", "o": "theory", "t": "The presentation lacks necessary conceptual clarity."}, {"i": 4, "phase": "post", "chg": null, "v": "positive", "o": "related_work", "t": "The claim of analyzing the joint effect is substantiated."}, {"i": 5, "phase": "post", "chg": null, "v": "positive", "o": "stats_metrics", "t": "The assumption is empirically grounded."}, {"i": 6, "phase": "post", "chg": null, "v": "positive", "o": "stats_metrics", "t": "The empirical comparison gap is closed."}, {"i": 7, "phase": "post", "chg": null, "v": "positive", "o": "theory", "t": "Conceptual clarity is improved."}, {"i": 8, "phase": "post", "chg": null, "v": "conditional", "o": "clarity", "t": "Satisfaction is contingent on final textual revision."}]}, {"key": "yc4e", "rating": 4, "units": [{"i": 0, "phase": "init", "chg": null, "v": "negative", "o": "novelty", "t": "The core contribution lacks sufficient novelty."}, {"i": 1, "phase": "init", "chg": null, "v": "negative", "o": "empirical_scope", "t": "The empirical value is weak."}, {"i": 2, "phase": "init", "chg": null, "v": "negative", "o": "stats_metrics", "t": "The empirical evaluation is insufficiently comprehensive."}, {"i": 3, "phase": "init", "chg": null, "v": "negative", "o": "robustness_sensitivity", "t": "The experimental setup lacks rigor."}, {"i": 4, "phase": "post", "chg": null, "v": "mixed", "o": "novelty", "t": "The novelty argument is partially addressed but remains debatable."}, {"i": 5, "phase": "post", "chg": null, "v": "positive", "o": "baselines_ablations", "t": "The empirical comparison landscape is significantly improved."}, {"i": 6, "phase": "post", "chg": null, "v": "positive", "o": "stats_metrics", "t": "The hyperparameter selection process is validated."}, {"i": 7, "phase": "post", "chg": null, "v": "positive", "o": "compute_cost", "t": "The efficiency argument is solid."}, {"i": 8, "phase": "post", "chg": null, "v": "negative", "o": "baselines_ablations", "t": "Uncertainty remains regarding the novelty threshold."}]}, {"key": "UTJo", "rating": 4, "units": [{"i": 0, "phase": "init", "chg": null, "v": "negative", "o": "empirical_scope", "t": "The empirical credibility of the small gains is uncertain."}, {"i": 1, "phase": "init", "chg": null, "v": "negative", "o": "empirical_scope", "t": "The practical feasibility is unverified."}, {"i": 2, "phase": "init", "chg": null, "v": "negative", "o": "empirical_scope", "t": "The robustness of the method in edge cases is unknown."}, {"i": 3, "phase": "post", "chg": null, "v": "positive", "o": "stats_metrics", "t": "The statistical credibility of the gains is improved."}, {"i": 4, "phase": "post", "chg": null, "v": "positive", "o": "compute_cost", "t": "The method is computationally efficient."}, {"i": 5, "phase": "post", "chg": null, "v": "positive", "o": "stats_metrics", "t": "The method's limitations are understood and manageable."}, {"i": 6, "phase": "post", "chg": null, "v": "negative", "o": "stats_metrics", "t": "Uncertainty remains regarding the value of small gains."}]}]}, {"scenario": "unanimity", "forum": "8oRi7CEj6A", "title": "Context Parametrization with Compositional Adapters", "decision": null, "comments": 6, "reviewers": [{"key": "Reviewer_WtSR", "rating": 6, "units": [{"i": 0, "phase": "init", "chg": null, "v": "positive", "o": "theory", "t": "The method's conceptual simplicity is a strength because treating context as composable adapters provides a practical alternative to long prompts."}, {"i": 1, "phase": "init", "chg": null, "v": "positive", "o": "baselines_ablations", "t": "The paper is correct but needs more boundary-mapping to be fully convincing."}, {"i": 2, "phase": "init", "chg": null, "v": "negative", "o": "compute_cost", "t": "The mechanism is not yet disambiguated regarding scaling limits."}, {"i": 3, "phase": "init", "chg": null, "v": "negative", "o": "empirical_scope", "t": "The robustness of the additive assumption is unverified."}, {"i": 4, "phase": "init", "chg": null, "v": "negative", "o": "compute_cost", "t": "The efficiency claims are incomplete due to inconsistent baseline reporting."}, {"i": 5, "phase": "init", "chg": null, "v": "negative", "o": "theory", "t": "The scaling properties are insufficiently characterized."}, {"i": 6, "phase": "init", "chg": null, "v": "negative", "o": "clarity", "t": "The theoretical formulation contains a confusing notational element."}, {"i": 7, "phase": "init", "chg": null, "v": "negative", "o": "robustness_sensitivity", "t": "The identifiability of summed adapters is unclear."}, {"i": 8, "phase": "init", "chg": null, "v": "negative", "o": "compute_cost", "t": "Resource usage claims lack quantitative backing."}, {"i": 9, "phase": "post", "chg": "strengthened", "v": "positive", "o": "compute_cost", "t": "The boundary-mapping gap is addressed by demonstrating clear patterns in block size variation and superior scaling compared to ICL."}, {"i": 10, "phase": "post", "chg": "strengthened", "v": "positive", "o": "empirical_scope", "t": "The fragility concern regarding additivity assumptions is mitigated by evidence of graceful degradation."}, {"i": 11, "phase": "post", "chg": "strengthened", "v": "positive", "o": "compute_cost", "t": "The efficiency comparison is now consistent and convincing."}, {"i": 12, "phase": "post", "chg": "strengthened", "v": "positive", "o": "clarity", "t": "The theoretical clarity issue is resolved."}, {"i": 13, "phase": "post", "chg": null, "v": "negative", "o": "theory", "t": "The potential stability issue at very large k remains unresolved."}, {"i": 14, "phase": "post", "chg": "strengthened", "v": "positive", "o": "compute_cost", "t": "The lack of quantitative backing for efficiency claims is resolved."}, {"i": 15, "phase": "post", "chg": "strengthened", "v": "positive", "o": "theory", "t": "The scaling characterization gap is filled."}]}, {"key": "Reviewer_jjPy", "rating": 6, "units": [{"i": 0, "phase": "init", "chg": null, "v": "positive", "o": "theory", "t": "The method has a solid theoretical foundation supporting its structure."}, {"i": 1, "phase": "init", "chg": null, "v": "positive", "o": "baselines_ablations", "t": "The necessity of each component is empirically validated."}, {"i": 2, "phase": "init", "chg": null, "v": "positive", "o": "theory", "t": "The method shows promising cross-backbone stability."}, {"i": 3, "phase": "init", "chg": null, "v": "negative", "o": "baselines_ablations", "t": "The validation is under-validated due to weak baselines."}, {"i": 4, "phase": "init", "chg": null, "v": "negative", "o": "compute_cost", "t": "Efficiency claims lack full cost transparency."}, {"i": 5, "phase": "init", "chg": null, "v": "negative", "o": "empirical_scope", "t": "Generalization to other task types is unproven."}, {"i": 6, "phase": "init", "chg": null, "v": "negative", "o": "robustness_sensitivity", "t": "Sensitivity to generator capacity is unexplored."}, {"i": 7, "phase": "post", "chg": "strengthened", "v": "positive", "o": "baselines_ablations", "t": "The method's superiority over stronger baselines is established."}, {"i": 8, "phase": "post", "chg": "strengthened", "v": "positive", "o": "compute_cost", "t": "Cost transparency is achieved through comprehensive reporting."}, {"i": 9, "phase": "post", "chg": "strengthened", "v": "positive", "o": "empirical_scope", "t": "Generalization to summarization is validated."}, {"i": 10, "phase": "post", "chg": "strengthened", "v": "mixed", "o": "compute_cost", "t": "Capacity sensitivity is clarified, though results are asserted without a table."}, {"i": 11, "phase": "post", "chg": null, "v": "negative", "o": "compute_cost", "t": "Efficiency generalization remains partially unverified."}]}, {"key": "Reviewer_UwJk", "rating": 6, "units": [{"i": 0, "phase": "init", "chg": null, "v": "positive", "o": "problem_framing", "t": "The method is promising and well-motivated."}, {"i": 1, "phase": "init", "chg": null, "v": "negative", "o": "method_design", "t": "The experimental design lacks necessary transparency for reproducibility."}, {"i": 2, "phase": "init", "chg": null, "v": "negative", "o": "empirical_scope", "t": "The experimental comparison may be confounded by unequal training data access."}, {"i": 3, "phase": "init", "chg": null, "v": "negative", "o": "compute_cost", "t": "Applicability to structured outputs is unverified."}, {"i": 4, "phase": "init", "chg": null, "v": "negative", "o": "empirical_scope", "t": "Compositionality with heterogeneous inputs is untested."}, {"i": 5, "phase": "post", "chg": "strengthened", "v": "positive", "o": "method_design", "t": "Transparency regarding training setup is improved."}, {"i": 6, "phase": "post", "chg": "strengthened", "v": "positive", "o": "baselines_ablations", "t": "The validity of reported gains is strengthened by fairer comparisons."}, {"i": 7, "phase": "post", "chg": "strengthened", "v": "positive", "o": "compute_cost", "t": "Applicability to summarization is validated."}, {"i": 8, "phase": "post", "chg": "strengthened", "v": "positive", "o": "empirical_scope", "t": "Compositionality with heterogeneous inputs is supported by experimental evidence."}, {"i": 9, "phase": "post", "chg": null, "v": "negative", "o": "stats_metrics", "t": "Statistical reliability of summarization results is questionable."}]}]}], "population": {"n_panels": 17848, "softening": 1006, "entrenchment": 1450, "reversal": 239, "split": 9329, "unanimity": 1127, "no_movement": 15289}, "move_units": {"strengthened": 3543, "clarified": 252, "weakened": 1053, "reversed": 283}}
