@misc{indiciaefee9a93bf2ed, title = {Learning Infinite-horizon Average-reward MDPs with Linear Function Approximation}, author = {Chen-Yu Wei and Mehdi Jafarnia-Jahromi and Haipeng Luo and Rahul Jain}, year = {2021}, url = {https://arxiv.org/abs/2007.11849}, note = {Source identifier: 2007.11849} }