@misc{indiciae976582fc7b30, title = {Posterior Sampling Reinforcement Learning with Gaussian Processes for Continuous Control: Sublinear Regret Bounds for Unbounded State Spaces}, author = {Hamish Flynn and Joe Watson and Ingmar Posner and Jan Peters}, year = {2026}, url = {https://arxiv.org/abs/2603.08287}, note = {Source identifier: 2603.08287} }