@misc{indiciaefe1d6e251c07, title = {SuperRL: Reinforcement Learning with Supervision to Boost Language Model Reasoning}, author = {Yihao Liu and Shuocheng Li and Lang Cao and Yuhang Xie and Mengyu Zhou and Haoyu Dong and Xiaojun Ma and Shi Han and Dongmei Zhang}, year = {2025}, url = {https://arxiv.org/abs/2506.01096}, note = {Source identifier: 2506.01096} }