diff sat/plugins/plugin_exp_lang_detect.py @ 2562:26edcf3a30eb

core, setup: huge cleaning: - moved directories from src and frontends/src to sat and sat_frontends, which is the recommanded naming convention - move twisted directory to root - removed all hacks from setup.py, and added missing dependencies, it is now clean - use https URL for website in setup.py - removed "Environment :: X11 Applications :: GTK", as wix is deprecated and removed - renamed sat.sh to sat and fixed its installation - added python_requires to specify Python version needed - replaced glib2reactor which use deprecated code by gtk3reactor sat can now be installed directly from virtualenv without using --system-site-packages anymore \o/
author Goffi <goffi@goffi.org>
date Mon, 02 Apr 2018 19:44:50 +0200
parents src/plugins/plugin_exp_lang_detect.py@0046283a285d
children 56f94936df1e
line wrap: on
line diff
--- /dev/null	Thu Jan 01 00:00:00 1970 +0000
+++ b/sat/plugins/plugin_exp_lang_detect.py	Mon Apr 02 19:44:50 2018 +0200
@@ -0,0 +1,91 @@
+#!/usr/bin/env python2
+# -*- coding: utf-8 -*-
+
+# SAT plugin to detect language (experimental)
+# Copyright (C) 2009-2018 Jérôme Poisson (goffi@goffi.org)
+
+# This program is free software: you can redistribute it and/or modify
+# it under the terms of the GNU Affero General Public License as published by
+# the Free Software Foundation, either version 3 of the License, or
+# (at your option) any later version.
+
+# This program is distributed in the hope that it will be useful,
+# but WITHOUT ANY WARRANTY; without even the implied warranty of
+# MERCHANTABILITY or FITNESS FOR A PARTICULAR PURPOSE.  See the
+# GNU Affero General Public License for more details.
+
+# You should have received a copy of the GNU Affero General Public License
+# along with this program.  If not, see <http://www.gnu.org/licenses/>.
+
+from sat.core.i18n import _, D_
+from sat.core.constants import Const as C
+from sat.core.log import getLogger
+log = getLogger(__name__)
+from sat.core import exceptions
+
+try:
+    from langid.langid import LanguageIdentifier, model
+except ImportError:
+    raise exceptions.MissingModule(u'Missing module langid, please download/install it with "pip install langid")')
+
+identifier = LanguageIdentifier.from_modelstring(model, norm_probs=False)
+
+
+PLUGIN_INFO = {
+    C.PI_NAME: "Language detection plugin",
+    C.PI_IMPORT_NAME: "EXP-LANG-DETECT",
+    C.PI_TYPE: "EXP",
+    C.PI_PROTOCOLS: [],
+    C.PI_DEPENDENCIES: [],
+    C.PI_MAIN: "LangDetect",
+    C.PI_HANDLER: "no",
+    C.PI_DESCRIPTION: _("""Detect and set message language when unknown""")
+}
+
+CATEGORY = D_(u"Misc")
+NAME = u"lang_detect"
+LABEL = D_(u"language detection")
+PARAMS = """
+    <params>
+    <individual>
+    <category name="{category_name}">
+        <param name="{name}" label="{label}" type="bool" value="true" />
+    </category>
+    </individual>
+    </params>
+    """.format(category_name=CATEGORY,
+               name=NAME,
+               label=_(LABEL),
+              )
+
+
+class LangDetect(object):
+
+    def __init__(self, host):
+        log.info(_(u"Language detection plugin initialization"))
+        self.host = host
+        host.memory.updateParams(PARAMS)
+        host.trigger.add("MessageReceived", self.MessageReceivedTrigger)
+        host.trigger.add("sendMessage", self.MessageSendTrigger)
+
+    def addLanguage(self, mess_data):
+        message = mess_data['message']
+        if len(message) == 1 and message.keys()[0] == '':
+            msg = message.values()[0]
+            lang = identifier.classify(msg)[0]
+            mess_data["message"] = {lang: msg}
+        return mess_data
+
+    def MessageReceivedTrigger(self, client, message_elt, post_treat):
+        """ Check if source is linked and repeat message, else do nothing  """
+
+        lang_detect = self.host.memory.getParamA(NAME, CATEGORY, profile_key=client.profile)
+        if lang_detect:
+            post_treat.addCallback(self.addLanguage)
+        return True
+
+    def MessageSendTrigger(self, client, data, pre_xml_treatments, post_xml_treatments):
+        lang_detect = self.host.memory.getParamA(NAME, CATEGORY, profile_key=client.profile)
+        if lang_detect:
+            self.addLanguage(data)
+        return True