@misc{indiciaed7c7ffaa1683, title = {S-GRPO: Unified Post-Training for Large Vision-Language Models}, author = {Yuming Yan and Kai Tang and Sihong Chen and Ke Xu and Dan Hu and Qun Yu and Pengfei Hu}, year = {2026}, url = {https://arxiv.org/abs/2604.16557}, note = {Source identifier: 2604.16557} }