@misc{indiciae25a7f3c3e1e6, title = {Confident, Calibrated, or Complicit: Safety Alignment and Ideological Bias in LLM Hate Speech Detection}, author = {Sanjeeevan Selvaganapathy and Mehwish Nasim}, year = {2026}, url = {https://arxiv.org/abs/2509.00673}, note = {Source identifier: 2509.00673} }