@misc{indiciaee032d062c0ef, title = {Causal Policy Learning in Reinforcement Learning: Backdoor-Adjusted Soft Actor-Critic}, author = {Thanh Vinh Vo and Young Lee and Haozhe Ma and Chien Lu and Tze-Yun Leong}, year = {2025}, url = {https://arxiv.org/abs/2506.05445}, note = {Source identifier: 2506.05445} }