@misc{indiciae94bb956a092e, title = {Tackling Heavy-Tailed Rewards in Reinforcement Learning with Function Approximation: Minimax Optimal and Instance-Dependent Regret Bounds}, author = {Jiayi Huang and Han Zhong and Liwei Wang and Lin F. Yang}, year = {2024}, url = {https://arxiv.org/abs/2306.06836}, note = {Source identifier: 2306.06836} }