@misc{indiciaeb7a9c6ca8abe, title = {Optimism in Reinforcement Learning and Kullback-Leibler Divergence}, author = {Sarah Filippi and Olivier Cappé and Aurélien Garivier}, year = {2010}, doi = {10.1109/allerton.2010.5706896}, url = {https://arxiv.org/abs/1004.5229}, note = {Source identifier: 1004.5229} }