Deprecated: The each() function is deprecated. This message will be suppressed on further calls in /home/zhenxiangba/zhenxiangba.com/public_html/phproxy-improved-master/index.php on line 456
JavaneseHonorifics (JavaneseHonorifics)
[go: Go Back, main page]

https://javanesehonorifics.github.io/\n
  • Paper: https://aclanthology.org/2025.acl-long.1296.pdf
  • \n
  • Repository: https://github.com/JavaneseHonorifics
  • \n
  • Venue: ACL 2025 Main Conference
  • \n\n

    Citation

    \n
    @inproceedings{farhansyah-etal-2025-language,\n    title = \"Do Language Models Understand Honorific Systems in {J}avanese?\",\n    author = \"Farhansyah, Mohammad Rifqi  and\n      Darmawan, Iwan  and\n      Kusumawardhana, Adryan  and\n      Winata, Genta Indra  and\n      Aji, Alham Fikri  and\n      Wijaya, Derry Tanti\",\n    editor = \"Che, Wanxiang  and\n      Nabende, Joyce  and\n      Shutova, Ekaterina  and\n      Pilehvar, Mohammad Taher\",\n    booktitle = \"Proceedings of the 63rd Annual Meeting of the Association for Computational Linguistics (Volume 1: Long Papers)\",\n    month = jul,\n    year = \"2025\",\n    address = \"Vienna, Austria\",\n    publisher = \"Association for Computational Linguistics\",\n    url = \"https://aclanthology.org/2025.acl-long.1296/\",\n    doi = \"10.18653/v1/2025.acl-long.1296\",\n    pages = \"26732--26754\",\n    ISBN = \"979-8-89176-251-0\",\n    abstract = \"The Javanese language features a complex system of honorifics that vary according to the social status of the speaker, listener, and referent. Despite its cultural and linguistic significance, there has been limited progress in developing a comprehensive corpus to capture these variations for natural language processing (NLP) tasks. In this paper, we present Unggah-Ungguh, a carefully curated dataset designed to encapsulate the nuances of Unggah-Ungguh Basa, the Javanese speech etiquette framework that dictates the choice of words and phrases based on social hierarchy and context. Using Unggah-Ungguh, we assess the ability of language models (LMs) to process various levels of Javanese honorifics through classification and machine translation tasks. To further evaluate cross-lingual LMs, we conduct machine translation experiments between Javanese (at specific honorific levels) and Indonesian. Additionally, we explore whether LMs can generate contextually appropriate Javanese honorifics in conversation tasks, where the honorific usage should align with the social role and contextual cues. Our findings indicate that current LMs struggle with most honorific levels, exhibiting a bias toward certain honorific tiers.\"\n}\n
    \n","classNames":"hf-sanitized hf-sanitized-BrRv5Z6saTQCxUGXiBQZq"},"users":[{"_id":"65af40481216d503271e7d95","avatarUrl":"https://cdn-avatars.huggingface.co/v1/production/uploads/65af40481216d503271e7d95/7BvAFiSHbL4rShMzdsLyd.png","isPro":false,"fullname":"Mohammad Rifqi Farhansyah","user":"rifqifarhansyah","type":"user"},{"_id":"5f5c4b20e56d546cd6233098","avatarUrl":"https://cdn-avatars.huggingface.co/v1/production/uploads/1637813888895-5f5c4b20e56d546cd6233098.jpeg","isPro":false,"fullname":"Genta Indra Winata","user":"gentaiscool","type":"user"},{"_id":"662084cf12f69ac918bb4a97","avatarUrl":"/avatars/01633a3b633d4e79f8995e2d1212e487.svg","isPro":false,"fullname":"Wijaya","user":"DerryTanti","type":"user"},{"_id":"665e9e32c19f7ccea08a46c9","avatarUrl":"/avatars/97fb1b043e9a81c4b70fcbfa729ea89a.svg","isPro":false,"fullname":"iwan d","user":"iwand","type":"user"},{"_id":"61a4dc053205e107691e0d82","avatarUrl":"https://cdn-avatars.huggingface.co/v1/production/uploads/61a4dc053205e107691e0d82/BESoEHlHYXstXudh6dOdT.jpeg","isPro":true,"fullname":"Alham Fikri Aji","user":"afaji","type":"user"}],"userCount":5,"collections":[],"datasets":[{"author":"JavaneseHonorifics","downloads":55,"gated":false,"id":"JavaneseHonorifics/Unggah-Ungguh","lastModified":"2026-01-21T05:16:53.000Z","datasetsServerInfo":{"viewer":"viewer","numRows":4184,"libraries":["datasets","pandas","mlcroissant","polars"],"formats":["csv"],"modalities":["tabular","text"]},"private":false,"repoType":"dataset","likes":1,"isLikedByUser":false,"isBenchmark":false}],"models":[{"author":"JavaneseHonorifics","authorData":{"_id":"68280a0137df496a1af39798","avatarUrl":"https://cdn-avatars.huggingface.co/v1/production/uploads/65af40481216d503271e7d95/Zc8R8HMemgGOKHvmTavYr.png","fullname":"JavaneseHonorifics","name":"JavaneseHonorifics","type":"org","isHf":false,"isHfAdmin":false,"isMod":false,"followerCount":5,"isUserFollowing":false},"downloads":7,"gated":false,"id":"JavaneseHonorifics/Unggah-Ungguh-Javanese-Distilbert-Classifier","availableInferenceProviders":[],"lastModified":"2026-01-21T05:16:19.000Z","likes":0,"pipeline_tag":"text-classification","private":false,"repoType":"model","isLikedByUser":false,"numParameters":66956548},{"author":"JavaneseHonorifics","authorData":{"_id":"68280a0137df496a1af39798","avatarUrl":"https://cdn-avatars.huggingface.co/v1/production/uploads/65af40481216d503271e7d95/Zc8R8HMemgGOKHvmTavYr.png","fullname":"JavaneseHonorifics","name":"JavaneseHonorifics","type":"org","isHf":false,"isHfAdmin":false,"isMod":false,"followerCount":5,"isUserFollowing":false},"downloads":7,"gated":false,"id":"JavaneseHonorifics/Unggah-Ungguh-Javanese-Bert-Classifier","availableInferenceProviders":[],"lastModified":"2026-01-21T05:16:00.000Z","likes":0,"pipeline_tag":"text-classification","private":false,"repoType":"model","isLikedByUser":false,"numParameters":109485316},{"author":"JavaneseHonorifics","authorData":{"_id":"68280a0137df496a1af39798","avatarUrl":"https://cdn-avatars.huggingface.co/v1/production/uploads/65af40481216d503271e7d95/Zc8R8HMemgGOKHvmTavYr.png","fullname":"JavaneseHonorifics","name":"JavaneseHonorifics","type":"org","isHf":false,"isHfAdmin":false,"isMod":false,"followerCount":5,"isUserFollowing":false},"downloads":8,"gated":false,"id":"JavaneseHonorifics/Unggah-Ungguh-Javanese-GPT2-Classifier","availableInferenceProviders":[],"lastModified":"2026-01-21T05:15:45.000Z","likes":0,"pipeline_tag":"text-classification","private":false,"repoType":"model","isLikedByUser":false,"numParameters":124442880},{"author":"JavaneseHonorifics","authorData":{"_id":"68280a0137df496a1af39798","avatarUrl":"https://cdn-avatars.huggingface.co/v1/production/uploads/65af40481216d503271e7d95/Zc8R8HMemgGOKHvmTavYr.png","fullname":"JavaneseHonorifics","name":"JavaneseHonorifics","type":"org","isHf":false,"isHfAdmin":false,"isMod":false,"followerCount":5,"isUserFollowing":false},"downloads":0,"gated":false,"id":"JavaneseHonorifics/Unggah-Ungguh-Javanese-LSTM-Classifier","availableInferenceProviders":[],"lastModified":"2026-01-21T05:15:31.000Z","likes":0,"pipeline_tag":"text-classification","private":false,"repoType":"model","isLikedByUser":false}],"paperPreviews":[],"spaces":[],"buckets":[],"numBuckets":0,"numDatasets":1,"numModels":4,"numSpaces":1,"lastOrgActivities":[{"time":"2026-01-27T14:40:03.193Z","user":"afaji","userAvatarUrl":"https://cdn-avatars.huggingface.co/v1/production/uploads/61a4dc053205e107691e0d82/BESoEHlHYXstXudh6dOdT.jpeg","type":"paper","paper":{"id":"2601.17277","title":"PingPong: A Natural Benchmark for Multi-Turn Code-Switching Dialogues","publishedAt":"2026-01-24T03:31:08.000Z","upvotes":6,"isUpvotedByUser":true}},{"time":"2026-01-27T14:39:54.347Z","user":"gentaiscool","userAvatarUrl":"https://cdn-avatars.huggingface.co/v1/production/uploads/1637813888895-5f5c4b20e56d546cd6233098.jpeg","type":"paper","paper":{"id":"2601.17277","title":"PingPong: A Natural Benchmark for Multi-Turn Code-Switching Dialogues","publishedAt":"2026-01-24T03:31:08.000Z","upvotes":6,"isUpvotedByUser":true}},{"time":"2026-01-27T09:42:31.965Z","user":"rifqifarhansyah","userAvatarUrl":"https://cdn-avatars.huggingface.co/v1/production/uploads/65af40481216d503271e7d95/7BvAFiSHbL4rShMzdsLyd.png","type":"paper","paper":{"id":"2601.17277","title":"PingPong: A Natural Benchmark for Multi-Turn Code-Switching Dialogues","publishedAt":"2026-01-24T03:31:08.000Z","upvotes":6,"isUpvotedByUser":true}}],"acceptLanguages":["*"],"canReadRepos":false,"canReadSpaces":false,"blogPosts":[],"currentRepoPage":0,"filters":{},"paperView":false}">

    AI & ML interests

    Low Resource Languages

    Recent Activity

    Javanese Honorifics (Unggah-Ungguh v1.0)

    The Javanese language, spoken by over 98 million people, features a distinctive honorific system known as Unggah-Ungguh Basa. We present UNGGAH-UNGGUH, a carefully curated dataset designed to encapsulate the nuances of Unggah-Ungguh Basa, the Javanese speech etiquette framework that dictates the choice of words and phrases based on social hierarchy and context.

    Citation

    @inproceedings{farhansyah-etal-2025-language,
        title = "Do Language Models Understand Honorific Systems in {J}avanese?",
        author = "Farhansyah, Mohammad Rifqi  and
          Darmawan, Iwan  and
          Kusumawardhana, Adryan  and
          Winata, Genta Indra  and
          Aji, Alham Fikri  and
          Wijaya, Derry Tanti",
        editor = "Che, Wanxiang  and
          Nabende, Joyce  and
          Shutova, Ekaterina  and
          Pilehvar, Mohammad Taher",
        booktitle = "Proceedings of the 63rd Annual Meeting of the Association for Computational Linguistics (Volume 1: Long Papers)",
        month = jul,
        year = "2025",
        address = "Vienna, Austria",
        publisher = "Association for Computational Linguistics",
        url = "https://aclanthology.org/2025.acl-long.1296/",
        doi = "10.18653/v1/2025.acl-long.1296",
        pages = "26732--26754",
        ISBN = "979-8-89176-251-0",
        abstract = "The Javanese language features a complex system of honorifics that vary according to the social status of the speaker, listener, and referent. Despite its cultural and linguistic significance, there has been limited progress in developing a comprehensive corpus to capture these variations for natural language processing (NLP) tasks. In this paper, we present Unggah-Ungguh, a carefully curated dataset designed to encapsulate the nuances of Unggah-Ungguh Basa, the Javanese speech etiquette framework that dictates the choice of words and phrases based on social hierarchy and context. Using Unggah-Ungguh, we assess the ability of language models (LMs) to process various levels of Javanese honorifics through classification and machine translation tasks. To further evaluate cross-lingual LMs, we conduct machine translation experiments between Javanese (at specific honorific levels) and Indonesian. Additionally, we explore whether LMs can generate contextually appropriate Javanese honorifics in conversation tasks, where the honorific usage should align with the social role and contextual cues. Our findings indicate that current LMs struggle with most honorific levels, exhibiting a bias toward certain honorific tiers."
    }