@misc{indiciae06b48b2421ab, title = {Sample-efficient Learning of Infinite-horizon Average-reward MDPs with General Function Approximation}, author = {Jianliang He and Han Zhong and Zhuoran Yang}, year = {2024}, url = {https://arxiv.org/abs/2404.12648}, note = {Source identifier: 2404.12648} }