@misc{indiciae3c7a41230fdb, title = {More Benefits of Being Distributional: Second-Order Bounds for Reinforcement Learning}, author = {Kaiwen Wang and Owen Oertell and Alekh Agarwal and Nathan Kallus and Wen Sun}, year = {2024}, url = {https://arxiv.org/abs/2402.07198}, note = {Source identifier: 2402.07198} }