@misc{indiciae496b952f9be0, title = {Robust Average-Reward Markov Decision Processes: Minimax-Optimal Learning via Plug-in Reductions}, author = {Yuepeng Yang and Yuxin Chen and Yuejie Chi}, year = {2026}, url = {https://arxiv.org/abs/2608.06545}, note = {Source identifier: 2608.06545} }