@misc{indiciaeea354d3b62a8, title = {No Prompt Left Behind: Exploiting Zero-Variance Prompts in LLM Reinforcement Learning via Entropy-Guided Advantage Shaping}, author = {Thanh-Long V. Le and Myeongho Jeon and Kim Vu and Viet Lai and Eunho Yang}, year = {2026}, url = {https://arxiv.org/abs/2509.21880}, note = {Source identifier: 2509.21880} }