@misc{indiciae32e0f26514be, title = {MoE-GRPO: Optimizing Mixture-of-Experts via Reinforcement Learning in Vision-Language Models}, author = {Dohwan Ko and Jinyoung Park and Seoung Choi and Sanghyeok Lee and Seohyun Lee and Hyunwoo J. Kim}, year = {2026}, url = {https://arxiv.org/abs/2603.24984}, note = {Source identifier: 2603.24984} }