@misc{indiciae8cbcec99a68c, title = {Is Vanilla Policy Gradient Overlooked? Analyzing Deep Reinforcement Learning for Hanabi}, author = {Bram Grooten and Jelle Wemmenhove and Maurice Poot and Jim Portegies}, year = {2022}, url = {https://arxiv.org/abs/2203.11656}, note = {Source identifier: 2203.11656} }