@misc{indiciaee1e16a71b533, title = {CE-GPPO: Coordinating Entropy via Gradient-Preserving Clipping Policy Optimization in Reinforcement Learning}, author = {Zhenpeng Su and Leiyu Pan and Minxuan Lv and Yuntao Li and Wenping Hu and Fuzheng Zhang and Kun Gai and Guorui Zhou}, year = {2026}, url = {https://arxiv.org/abs/2509.20712}, note = {Source identifier: 2509.20712} }