@misc{indiciae875c98d7b623, title = {When Implausible Tokens Get Reinforced: Tail-Aware Credit Calibration for LLM Reinforcement Learning}, author = {Xiuyi Lou and Zicheng Xu and Yu-Neng Chuang and Hoang Anh Duy Le and Zhaozhuo Xu and Guanchu Wang and Vladimir Braverman}, year = {2026}, url = {https://arxiv.org/abs/2607.07976}, note = {Source identifier: 2607.07976} }