@misc{indiciaee62d58234c11, title = {Improving Neuron-level Interpretability with White-box Language Models}, author = {Hao Bai and Yi Ma}, year = {2025}, url = {https://arxiv.org/abs/2410.16443}, note = {Source identifier: 2410.16443} }