@misc{indiciae160685b3238d, title = {On the effectiveness of reward functions in reinforcement learning for confidence calibration of large language models}, author = {Chee Heng Tan and Zhuoyi Lin and Mehul Motani and Wee Sun Lee}, year = {2026}, url = {https://arxiv.org/abs/2607.04332}, note = {Source identifier: 2607.04332} }