@misc{indiciae006cf75d146d, title = {Rethinking Ratio-Based Trust Regions for Policy Optimization in Multi-Agent Reinforcement Learning}, author = {Chulabhaya Wijesundara and Andrea Baisero and Zhongheng Li and Gregory Castañón and Alan Carlin and Christopher Amato}, year = {2026}, url = {https://arxiv.org/abs/2605.09212}, note = {Source identifier: 2605.09212} }