@misc{indiciae9e0de4075055, title = {Bellman Unbiasedness: Toward Provably Efficient Distributional Reinforcement Learning with General Value Function Approximation}, author = {Taehyun Cho and Seungyub Han and Seokhun Ju and Dohyeong Kim and Kyungjae Lee and Jungwoo Lee}, year = {2025}, url = {https://arxiv.org/abs/2407.21260}, note = {Source identifier: 2407.21260} }