@misc{indiciaed51643b35c1c, title = {E-GRPO: High Entropy Steps Drive Effective Reinforcement Learning for Flow Models}, author = {Shengjun Zhang and Zhang Zhang and Chensheng Dai and Yueqi Duan}, year = {2026}, url = {https://arxiv.org/abs/2601.00423}, note = {Source identifier: 2601.00423} }