@misc{indiciaea576f9be9ea1, title = {Decoupling Reasoning and Confidence: Resurrecting Calibration in Reinforcement Learning from Verifiable Rewards}, author = {Zhengzhao Ma and Xueru Wen and Boxi Cao and Yaojie Lu and Hongyu Lin and Jinglin Yang and Min He and Xianpei Han and Le Sun}, year = {2026}, url = {https://arxiv.org/abs/2603.09117}, note = {Source identifier: 2603.09117} }