This is a sentence-transformers model finetuned from rasyosef/roberta-base-amharic. It maps sentences & paragraphs to a 768-dimensional dense vector space and can be used for semantic textual similarity, semantic search, paraphrase mining, text classification, clustering, and more.
1@inproceedings{alemneh2026amharicir,
2 title = {The Multilingual Curse at the Retrieval Layer: Evidence from Amharic},
3 author = {Alemneh, Yosef Worku and Mekonnen, Kidist Amde and de Rijke, Maarten},
4 booktitle = {Proceedings of the 1st Workshop on Multilinguality in the Era of Large Language Models (MeLLM), ACL 2026},
5 year = {2026},
6}