@misc{indiciae258bf0659bc6, title = {From n-gram to Attention: How Model Architectures Learn and Propagate Bias in Language Modeling}, author = {Mohsinul Kabir and Tasfia Tahsin and Sophia Ananiadou}, year = {2025}, doi = {10.18653/v1/2025.findings-emnlp.1003}, url = {https://arxiv.org/abs/2505.12381}, note = {Source identifier: 2505.12381} }