@misc{indiciaedf6bfc036cb1, title = {Efficient Exploration in Average-Reward Constrained Reinforcement Learning: Achieving Near-Optimal Regret With Posterior Sampling}, author = {Danil Provodin and Maurits Kaptein and Mykola Pechenizkiy}, year = {2024}, url = {https://arxiv.org/abs/2405.19017}, note = {Source identifier: 2405.19017} }