@misc{indiciae700a8c94d9db, title = {Adaptive Negative Reinforcement for LLM Reasoning:Dynamically Balancing Correction and Diversity in RLVR}, author = {Yash Ingle and Jaival Chauhan and Ankit Yadav and Sudhakar Mishra}, year = {2026}, url = {https://arxiv.org/abs/2605.07137}, note = {Source identifier: 2605.07137} }