@misc{indiciaeb443498f2ad8, title = {Why is Posterior Sampling Better than Optimism for Reinforcement Learning?}, author = {Ian Osband and Benjamin Van Roy}, year = {2017}, url = {https://arxiv.org/abs/1607.00215}, note = {Source identifier: 1607.00215} }