@misc{indiciae49e1210b4bc9, title = {Entropy-Gated Selective Policy Optimization:Token-Level Gradient Allocation for Hybrid Training of Large Language Models}, author = {Yuelin Hu and Zhengxue Cheng and Wei Liu and Li Song}, year = {2026}, url = {https://arxiv.org/abs/2602.03309}, note = {Source identifier: 2602.03309} }