@misc{indiciae82385ee2b43b, title = {LongR: Unleashing Long-Context Reasoning via Reinforcement Learning with Dense Utility Rewards}, author = {Bowen Ping and Zijun Chen and Yiyao Yu and Tingfeng Hui and Junchi Yan and Baobao Chang}, year = {2026}, url = {https://arxiv.org/abs/2602.05758}, note = {Source identifier: 2602.05758} }