@inproceedings{sarkar-etal-2025-zero,
title = "Zero-Shot Multi-Label Classification of {B}angla Documents: Large Decoders Vs. Classic Encoders",
author = "Sarkar, Souvika and
Hasan, Md Najib and
Karmaker, Santu",
editor = "Alam, Firoj and
Kar, Sudipta and
Chowdhury, Shammur Absar and
Hassan, Naeemul and
Prince, Enamul Hoque and
Tasnim, Mohiuddin and
Rony, Md Rashad Al Hasan and
Rahman, Md Tahmid Rahman",
booktitle = "Proceedings of the Second Workshop on Bangla Language Processing (BLP-2025)",
month = dec,
year = "2025",
address = "Mumbai, India",
publisher = "Association for Computational Linguistics",
url = "https://preview.aclanthology.org/ingest-ijcnlp-aacl/2025.banglalp-1.7/",
pages = "91--106",
ISBN = "979-8-89176-314-2",
abstract = "Bangla, a language spoken by over 300 million native speakers and ranked as the sixth most spoken language worldwide, presents unique challenges in natural language processing (NLP) due to its complex morphological characteristics and limited resources. Although recent large-decoder-based LLMs, such as GPT, LLaMA, and DeepSeek, have demonstrated excellent performance across many NLP tasks, their effectiveness in Bangla remains largely unexplored. In this paper, we establish the first benchmark comparing large decoder-based LLMs with classic encoder-based models for the Zero-Shot Multi-Label Classification (Zero-Shot-MLC) task in Bangla. Our evaluation of 32 state-of-the-art models reveals that existing so-called powerful encoders and decoders still struggle to achieve high accuracy on the Bangla Zero-Shot-MLC task, suggesting a need for more research and resources for Bangla NLP."
}Markdown (Informal)
[Zero-Shot Multi-Label Classification of Bangla Documents: Large Decoders Vs. Classic Encoders](https://preview.aclanthology.org/ingest-ijcnlp-aacl/2025.banglalp-1.7/) (Sarkar et al., BanglaLP 2025)
ACL