@misc{indiciaefb0bbadb6b23, title = {Efficient Q-Learning and Actor-Critic Methods for Robust Average-Reward Reinforcement Learning}, author = {Yang Xu and Swetha Ganesh and Vaneet Aggarwal}, year = {2026}, url = {https://arxiv.org/abs/2506.07040}, note = {Source identifier: 2506.07040} }