@misc{indiciae6e74cfc94894, title = {Action Gaps and Advantages in Continuous-Time Distributional Reinforcement Learning}, author = {Harley Wiltzer and Marc G. Bellemare and David Meger and Patrick Shafto and Yash Jhaveri}, year = {2024}, url = {https://arxiv.org/abs/2410.11022}, note = {Source identifier: 2410.11022} }