@misc{indiciae8042baa2ce3c, title = {Tighter Problem-Dependent Regret Bounds in Reinforcement Learning without Domain Knowledge using Value Function Bounds}, author = {Andrea Zanette and Emma Brunskill}, year = {2019}, url = {https://arxiv.org/abs/1901.00210}, note = {Source identifier: 1901.00210} }