@misc{indiciae7932b284a7a3, title = {Non-Linear Reinforcement Learning in Large Action Spaces: Structural Conditions and Sample-efficiency of Posterior Sampling}, author = {Alekh Agarwal and Tong Zhang}, year = {2024}, url = {https://arxiv.org/abs/2203.08248}, note = {Source identifier: 2203.08248} }