@misc{indiciae847e8cc2eb17, title = {Rewarding Doubt: A Reinforcement Learning Approach to Calibrated Confidence Expression of Large Language Models}, author = {David Bani-Harouni and Chantal Pellegrini and Paul Stangel and Ege Özsoy and Kamilia Zaripova and Nassir Navab and Matthias Keicher}, year = {2026}, url = {https://arxiv.org/abs/2503.02623}, note = {Source identifier: 2503.02623} }