@misc{indiciae4c449250e663, title = {Nearly Minimax Optimal Regret for Learning Infinite-horizon Average-reward MDPs with Linear Function Approximation}, author = {Yue Wu and Dongruo Zhou and Quanquan Gu}, year = {2022}, url = {https://arxiv.org/abs/2102.07301}, note = {Source identifier: 2102.07301} }