@misc{indiciaebad6f10a9d71, title = {Reinforcement Learning for Exponential Utility: Algorithms and Convergence in Discounted MDPs}, author = {Gugan Thoppe and L. A. Prashanth and Ankur Naskar and Sanjay Bhat}, year = {2026}, url = {https://arxiv.org/abs/2605.08053}, note = {Source identifier: 2605.08053} }