@misc{indiciae14e7d1523348, title = {Strategy Masking: A Method for Guardrails in Value-based Reinforcement Learning Agents}, author = {Jonathan Keane and Sam Keyser and Jeremy Kedziora}, year = {2025}, url = {https://arxiv.org/abs/2501.05501}, note = {Source identifier: 2501.05501} }