Mercurial > libervia-backend
view src/plugins/plugin_exp_lang_detect.py @ 2013:b536dd121da1
backend (memory), frontends: improved history filtering:
a "filters" dictionnary is now use to filter, it can have, for now, filtering on:
- "body": filter only on the body (equivalent to former "search" parameter, but not case sensitive)
- "search": fitler on body + source resource
- "types": allowed types
- "not_types": forbidden types
primitivus now do searching using "search", i.e. source resource is now taken into account (and search is now case insensitive)
author | Goffi <goffi@goffi.org> |
---|---|
date | Mon, 18 Jul 2016 00:52:02 +0200 |
parents | d95a6d553bec |
children | 1d3f73e065e1 |
line wrap: on
line source
#!/usr/bin/env python2 # -*- coding: utf-8 -*- # SAT plugin to detect language (experimental) # Copyright (C) 2009-2016 Jérôme Poisson (goffi@goffi.org) # This program is free software: you can redistribute it and/or modify # it under the terms of the GNU Affero General Public License as published by # the Free Software Foundation, either version 3 of the License, or # (at your option) any later version. # This program is distributed in the hope that it will be useful, # but WITHOUT ANY WARRANTY; without even the implied warranty of # MERCHANTABILITY or FITNESS FOR A PARTICULAR PURPOSE. See the # GNU Affero General Public License for more details. # You should have received a copy of the GNU Affero General Public License # along with this program. If not, see <http://www.gnu.org/licenses/>. from sat.core.i18n import _, D_ from sat.core.log import getLogger log = getLogger(__name__) from sat.core import exceptions try: from langid.langid import LanguageIdentifier, model except ImportError: raise exceptions.MissingModule(u'Missing module langid, please download/install it with "pip install langid")') identifier = LanguageIdentifier.from_modelstring(model, norm_probs=False) PLUGIN_INFO = { "name": "Language detection plugin", "import_name": "EXP-LANG-DETECT", "type": "EXP", "protocols": [], "dependencies": [], "main": "LangDetect", "handler": "no", "description": _("""Detect and set message language when unknown""") } CATEGORY = D_(u"Misc") NAME = u"lang_detect" LABEL = D_(u"language detection") PARAMS = """ <params> <individual> <category name="{category_name}"> <param name="{name}" label="{label}" type="bool" value="true" /> </category> </individual> </params> """.format(category_name=CATEGORY, name=NAME, label=_(LABEL), ) class LangDetect(object): def __init__(self, host): log.info(_(u"Language detection plugin initialization")) self.host = host host.memory.updateParams(PARAMS) host.trigger.add("MessageReceived", self.MessageReceivedTrigger) host.trigger.add("messageSend", self.MessageSendTrigger) def addLanguage(self, mess_data): message = mess_data['message'] if len(message) == 1 and message.keys()[0] == '': msg = message.values()[0] lang = identifier.classify(msg)[0] mess_data["message"] = {lang: msg} return mess_data def MessageReceivedTrigger(self, client, message_elt, post_treat): """ Check if source is linked and repeat message, else do nothing """ lang_detect = self.host.memory.getParamA(NAME, CATEGORY, profile_key=client.profile) if lang_detect: post_treat.addCallback(self.addLanguage) return True def MessageSendTrigger(self, client, data, pre_xml_treatments, post_xml_treatments): lang_detect = self.host.memory.getParamA(NAME, CATEGORY, profile_key=client.profile) if lang_detect: self.addLanguage(data) return True