first commit

2020-01-02 22:20:31 +07:00
commit 10eb3340ad
5753 changed files with 631345 additions and 0 deletions
--- a/administrator/components/com_finder/helpers/indexer/stemmer/fr.php
+++ b/administrator/components/com_finder/helpers/indexer/stemmer/fr.php
@@ -0,0 +1,264 @@
+<?php
+/**
+ * @package     Joomla.Administrator
+ * @subpackage  com_finder
+ *
+ * @copyright   Copyright (C) 2005 - 2013 Open Source Matters, Inc. All rights reserved.
+ * @license     GNU General Public License version 2 or later; see LICENSE
+ */
+
+defined('_JEXEC') or die;
+
+JLoader::register('FinderIndexerStemmer', dirname(__DIR__) . '/stemmer.php');
+
+/**
+ * French stemmer class for Smart Search indexer.
+ *
+ * First contributed by Eric Sanou (bobotche@hotmail.fr)
+ * This class is inspired in  Alexis Ulrich's French stemmer code (http://alx2002.free.fr)
+ *
+ * @package     Joomla.Administrator
+ * @subpackage  com_finder
+ * @since       3.0
+ */
+class FinderIndexerStemmerFr extends FinderIndexerStemmer
+{
+	/**
+	 * Stemming rules.
+	 *
+	 * @var    array
+	 * @since  3.0
+	 */
+	private static $_stemRules = null;
+
+	/**
+	 * Method to stem a token and return the root.
+	 *
+	 * @param   string  $token  The token to stem.
+	 * @param   string  $lang   The language of the token.
+	 *
+	 * @return  string  The root token.
+	 *
+	 * @since   3.0
+	 */
+	public function stem($token, $lang)
+	{
+		// Check if the token is long enough to merit stemming.
+		if (strlen($token) <= 2)
+		{
+			return $token;
+		}
+
+		// Check if the language is French or All.
+		if ($lang !== 'fr' && $lang != '*')
+		{
+			return $token;
+		}
+
+		// Stem the token if it is not in the cache.
+		if (!isset($this->cache[$lang][$token]))
+		{
+			// Stem the token.
+			$result = static::_getStem($token);
+
+			// Add the token to the cache.
+			$this->cache[$lang][$token] = $result;
+		}
+
+		return $this->cache[$lang][$token];
+	}
+
+	/**
+	 * French stemmer rules variables.
+	 *
+	 * @return  array  The rules
+	 *
+	 * @since   3.0
+	 */
+	protected static function getStemRules()
+	{
+		if (static::$_stemRules)
+		{
+			return static::$_stemRules;
+		}
+
+		$vars = array();
+
+		// French accented letters in ISO-8859-1 encoding
+		$vars['accents'] = chr(224) . chr(226) . chr(232) . chr(233) . chr(234) . chr(235) . chr(238) . chr(239) . chr(244) . chr(251) . chr(249) . chr(231);
+
+		// The rule patterns include all accented words for french language
+		$vars['rule_pattern'] = "/^([a-z" . $vars['accents'] . "]*)(\*){0,1}(\d)([a-z" . $vars['accents'] . "]*)([.|>])/";
+
+		// French vowels (including y) in ISO-8859-1 encoding
+		$vars['vowels'] = chr(97) . chr(224) . chr(226) . chr(101) . chr(232) . chr(233) . chr(234) . chr(235) . chr(105) . chr(238) . chr(239) . chr(111) . chr(244) . chr(117) . chr(251) . chr(249) . chr(121);
+
+		// The French rules in ISO-8859-1 encoding
+		$vars['rules'] = array(
+			'esre1>', 'esio1>', 'siol1.', 'siof0.', 'sioe0.', 'sio3>', 'st1>', 'sf1>', 'sle1>', 'slo1>', 's' . chr(233) . '1>', chr(233) . 'tuae5.',
+			chr(233) . 'tuae2.', 'tnia0.', 'tniv1.', 'tni3>', 'suor1.', 'suo0.', 'sdrail5.', 'sdrai4.', 'er' . chr(232) . 'i1>', 'sesue3x>',
+			'esuey5i.', 'esue2x>', 'se1>', 'er' . chr(232) . 'g3.', 'eca1>', 'esiah0.', 'esi1>', 'siss2.', 'sir2>', 'sit2>', 'egan' . chr(233) . '1.',
+			'egalli6>', 'egass1.', 'egas0.', 'egat3.', 'ega3>', 'ette4>', 'ett2>', 'etio1.', 'tio' . chr(231) . '4c.', 'tio0.', 'et1>', 'eb1>',
+			'snia1>', 'eniatnau8>', 'eniatn4.', 'enia1>', 'niatnio3.', 'niatg3.', 'e' . chr(233) . '1>', chr(233) . 'hcat1.', chr(233) . 'hca4.',
+			chr(233) . 'tila5>', chr(233) . 'tici5.', chr(233) . 'tir1.', chr(233) . 'ti3>', chr(233) . 'gan1.', chr(233) . 'ga3>',
+			chr(233) . 'tehc1.', chr(233) . 'te3>', chr(233) . 'it0.', chr(233) . '1>', 'eire4.', 'eirue5.', 'eio1.', 'eia1.', 'ei1>', 'eng1.',
+			'xuaessi7.', 'xuae1>', 'uaes0.', 'uae3.', 'xuave2l.', 'xuav2li>', 'xua3la>', 'ela1>', 'lart2.', 'lani2>', 'la' . chr(233) . '2>',
+			'siay4i.', 'siassia7.', 'siarv1*.', 'sia1>', 'tneiayo6i.', 'tneiay6i.', 'tneiassia9.', 'tneiareio7.', 'tneia5>', 'tneia4>', 'tiario4.',
+			'tiarim3.', 'tiaria3.', 'tiaris3.', 'tiari5.', 'tiarve6>', 'tiare5>', 'iare4>', 'are3>', 'tiay4i.', 'tia3>', 'tnay4i.',
+			'em' . chr(232) . 'iu5>', 'em' . chr(232) . 'i4>', 'tnaun3.', 'tnauqo3.', 'tnau4>', 'tnaf0.', 'tnat' . chr(233) . '2>', 'tna3>', 'tno3>',
+			'zeiy4i.', 'zey3i.', 'zeire5>', 'zeird4.', 'zeirio4.', 'ze2>', 'ssiab0.', 'ssia4.', 'ssi3.', 'tnemma6>', 'tnemesuey9i.', 'tnemesue8>',
+			'tnemevi7.', 'tnemessia5.', 'tnemessi8.', 'tneme5>', 'tnemia4.', 'tnem' . chr(233) . '5>', 'el2l>', 'lle3le>', 'let' . chr(244) . '0.',
+			'lepp0.', 'le2>', 'srei1>', 'reit3.', 'reila2.', 'rei3>', 'ert' . chr(226) . 'e5.', 'ert' . chr(226) . chr(233) . '1.',
+			'ert' . chr(226) . '4.', 'drai4.', 'erdro0.', 'erute5.', 'ruta0.', 'eruta1.', 'erutiov1.', 'erub3.', 'eruh3.', 'erul3.', 'er2r>', 'nn1>',
+			'r' . chr(232) . 'i3.', 'srev0.', 'sr1>', 'rid2>', 're2>', 'xuei4.', 'esuei5.', 'lbati3.', 'lba3>', 'rueis0.', 'ruehcn4.', 'ecirta6.',
+			'ruetai6.', 'rueta5.', 'rueir0.', 'rue3>', 'esseti6.', 'essere6>', 'esserd1.', 'esse4>', 'essiab1.', 'essia5.', 'essio1.', 'essi4.',
+			'essal4.', 'essa1>', 'ssab1.', 'essurp1.', 'essu4.', 'essi1.', 'ssor1.', 'essor2.', 'esso1>', 'ess2>', 'tio3.', 'r' . chr(232) . 's2re.',
+			'r' . chr(232) . '0e.', 'esn1.', 'eu1>', 'sua0.', 'su1>', 'utt1>', 'tu' . chr(231) . '3c.', 'u' . chr(231) . '2c.', 'ur1.', 'ehcn2>',
+			'ehcu1>', 'snorr3.', 'snoru3.', 'snorua3.', 'snorv3.', 'snorio4.', 'snori5.', 'snore5>', 'snortt4>', 'snort' . chr(238) . 'a7.', 'snort3.',
+			'snor4.', 'snossi6.', 'snoire6.', 'snoird5.', 'snoitai7.', 'snoita6.', 'snoits1>', 'noits0.', 'snoi4>', 'noitaci7>', 'noitai6.', 'noita5.',
+			'noitu4.', 'noi3>', 'snoya0.', 'snoy4i.', 'sno' . chr(231) . 'a1.', 'sno' . chr(231) . 'r1.', 'snoe4.', 'snosiar1>', 'snola1.', 'sno3>',
+			'sno1>', 'noll2.', 'tnennei4.', 'ennei2>', 'snei1>', 'sne' . chr(233) . '1>', 'enne' . chr(233) . '5e.', 'ne' . chr(233) . '3e.', 'neic0.',
+			'neiv0.', 'nei3.', 'sc1.', 'sd1.', 'sg1.', 'sni1.', 'tiu0.', 'ti2.', 'sp1>', 'sna1>', 'sue1.', 'enn2>', 'nong2.', 'noss2.', 'rioe4.',
+			'riot0.', 'riorc1.', 'riovec5.', 'rio3.', 'ric2.', 'ril2.', 'tnerim3.', 'tneris3>', 'tneri5.', 't' . chr(238) . 'a3.', 'riss2.',
+			't' . chr(238) . '2.', 't' . chr(226) . '2>', 'ario2.', 'arim1.', 'ara1.', 'aris1.', 'ari3.', 'art1>', 'ardn2.', 'arr1.', 'arua1.',
+			'aro1.', 'arv1.', 'aru1.', 'ar2.', 'rd1.', 'ud1.', 'ul1.', 'ini1.', 'rin2.', 'tnessiab3.', 'tnessia7.', 'tnessi6.', 'tnessni4.', 'sini2.',
+			'sl1.', 'iard3.', 'iario3.', 'ia2>', 'io0.', 'iule2.', 'i1>', 'sid2.', 'sic2.', 'esoi4.', 'ed1.', 'ai2>', 'a1>', 'adr1.',
+			'tner' . chr(232) . '5>', 'evir1.', 'evio4>', 'evi3.', 'fita4.', 'fi2>', 'enie1.', 'sare4>', 'sari4>', 'sard3.', 'sart2>', 'sa2.',
+			'tnessa6>', 'tnessu6>', 'tnegna3.', 'tnegi3.', 'tneg0.', 'tneru5>', 'tnemg0.', 'tnerni4.', 'tneiv1.', 'tne3>', 'une1.', 'en1>', 'nitn2.',
+			'ecnay5i.', 'ecnal1.', 'ecna4.', 'ec1>', 'nn1.', 'rit2>', 'rut2>', 'rud2.', 'ugn1>', 'eg1>', 'tuo0.', 'tul2>', 't' . chr(251) . '2>',
+			'ev1>', 'v' . chr(232) . '2ve>', 'rtt1>', 'emissi6.', 'em1.', 'ehc1.', 'c' . chr(233) . 'i2c' . chr(232) . '.', 'libi2l.', 'llie1.',
+			'liei4i.', 'xuev1.', 'xuey4i.', 'xueni5>', 'xuell4.', 'xuere5.', 'xue3>', 'rb' . chr(233) . '3rb' . chr(232) . '.', 'tur2.',
+			'rir' . chr(233) . '4re.', 'rir2.', 'c' . chr(226) . '2ca.', 'snu1.', 'rt' . chr(238) . 'a4.', 'long2.', 'vec2.', chr(231) . '1c>',
+			'ssilp3.', 'silp2.', 't' . chr(232) . 'hc2te.', 'n' . chr(232) . 'm2ne.', 'llepp1.', 'tan2.', 'rv' . chr(232) . '3rve.',
+			'rv' . chr(233) . '3rve.', 'r' . chr(232) . '2re.', 'r' . chr(233) . '2re.', 't' . chr(232) . '2te.', 't' . chr(233) . '2te.', 'epp1.',
+			'eya2i.', 'ya1i.', 'yo1i.', 'esu1.', 'ugi1.', 'tt1.', 'end0.'
+		);
+
+		static::$_stemRules = $vars;
+
+		return static::$_stemRules;
+	}
+
+	/**
+	 * Returns the number of the first rule from the rule number
+	 * that can be applied to the given reversed input.
+	 * returns -1 if no rule can be applied, ie the stem has been found
+	 *
+	 * @param   string   $reversed_input  The input to check in reversed order
+	 * @param   integer  $rule_number     The rule number to check
+	 *
+	 * @return  integer  Number of the first rule
+	 *
+	 * @since   3.0
+	 */
+	private static function _getFirstRule($reversed_input, $rule_number)
+	{
+		$vars = static::getStemRules();
+
+		$nb_rules = count($vars['rules']);
+
+		for ($i = $rule_number; $i < $nb_rules; $i++)
+		{
+			// Gets the letters from the current rule
+			$rule = $vars['rules'][$i];
+			$rule = preg_replace($vars['rule_pattern'], "\\1", $rule);
+
+			if (strncasecmp(utf8_decode($rule), $reversed_input, strlen(utf8_decode($rule))) == 0)
+			{
+				return $i;
+			}
+		}
+
+		return -1;
+	}
+
+	/**
+	 * Check the acceptability of a stem for French language
+	 *
+	 * @param   string  $reversed_stem  The stem to check in reverse form
+	 *
+	 * @return  boolean  True if stem is acceptable
+	 *
+	 * @since   3.0
+	 */
+	private static function _check($reversed_stem)
+	{
+		$vars = static::getStemRules();
+
+		if (preg_match('/[' . $vars['vowels'] . ']$/', utf8_encode($reversed_stem)))
+		{
+			// If the form starts with a vowel then at least two letters must remain after stemming (e.g.: "etaient" --> "et")
+			return (strlen($reversed_stem) > 2);
+		}
+		else
+		{
+			// If the reversed stem starts with a consonant then at least two letters must remain after stemming
+			if (strlen($reversed_stem) <= 2)
+			{
+				return false;
+			}
+
+			// And at least one of these must be a vowel or "y"
+			return (preg_match('/[' . $vars['vowels'] . ']/', utf8_encode($reversed_stem)));
+		}
+	}
+
+	/**
+	 * Paice/Husk stemmer which returns a stem for the given $input
+	 *
+	 * @param   string  $input  The word for which we want the stem in UTF-8
+	 *
+	 * @return  string  The stem
+	 *
+	 * @since   3.0
+	 */
+	private static function _getStem($input)
+	{
+		$vars = static::getStemRules();
+
+		$intact = true;
+		$reversed_input = strrev(utf8_decode($input));
+		$rule_number = 0;
+
+		// This loop goes through the rules' array until it finds an ending one (ending by '.') or the last one ('end0.')
+		while (true)
+		{
+			$rule_number = static::_getFirstRule($reversed_input, $rule_number);
+
+			if ($rule_number == -1)
+			{
+				// No other rule can be applied => the stem has been found
+				break;
+			}
+			$rule = $vars['rules'][$rule_number];
+			preg_match($vars['rule_pattern'], $rule, $matches);
+
+			if (($matches[2] != '*') || ($intact))
+			{
+				$reversed_stem = utf8_decode($matches[4]) . substr($reversed_input, $matches[3], strlen($reversed_input) - $matches[3]);
+
+				if (self::_check($reversed_stem))
+				{
+					$reversed_input = $reversed_stem;
+
+					if ($matches[5] == '.')
+					{
+						break;
+					}
+				}
+				else
+				{
+					// Go to another rule
+					$rule_number++;
+				}
+			}
+			else
+			{
+				// Go to another rule
+				$rule_number++;
+			}
+		}
+
+		return utf8_encode(strrev($reversed_input));
+	}
+}
--- a/administrator/components/com_finder/helpers/indexer/stemmer/index.html
+++ b/administrator/components/com_finder/helpers/indexer/stemmer/index.html
@@ -0,0 +1 @@
+<!DOCTYPE html><title></title>
--- a/administrator/components/com_finder/helpers/indexer/stemmer/porter_en.php
+++ b/administrator/components/com_finder/helpers/indexer/stemmer/porter_en.php
@@ -0,0 +1,448 @@
+<?php
+/**
+ * @package     Joomla.Administrator
+ * @subpackage  com_finder
+ *
+ * @copyright   Copyright (C) 2005 - 2013 Open Source Matters, Inc. All rights reserved.
+ * @license     GNU General Public License version 2 or later; see LICENSE
+ */
+
+defined('_JEXEC') or die;
+
+JLoader::register('FinderIndexerStemmer', dirname(__DIR__) . '/stemmer.php');
+
+/**
+ * Porter English stemmer class for the Finder indexer package.
+ *
+ * This class was adapted from one written by Richard Heyes.
+ * See copyright and link information above.
+ *
+ * @package     Joomla.Administrator
+ * @subpackage  com_finder
+ * @since       2.5
+ */
+class FinderIndexerStemmerPorter_En extends FinderIndexerStemmer
+{
+	/**
+	 * Regex for matching a consonant.
+	 *
+	 * @var    string
+	 * @since  2.5
+	 */
+	private static $_regex_consonant = '(?:[bcdfghjklmnpqrstvwxz]|(?<=[aeiou])y|^y)';
+
+	/**
+	 * Regex for matching a vowel
+	 *
+	 * @var    string
+	 * @since  2.5
+	 */
+	private static $_regex_vowel = '(?:[aeiou]|(?<![aeiou])y)';
+
+	/**
+	 * Method to stem a token and return the root.
+	 *
+	 * @param   string  $token  The token to stem.
+	 * @param   string  $lang   The language of the token.
+	 *
+	 * @return  string  The root token.
+	 *
+	 * @since   2.5
+	 */
+	public function stem($token, $lang)
+	{
+		// Check if the token is long enough to merit stemming.
+		if (strlen($token) <= 2)
+		{
+			return $token;
+		}
+
+		// Check if the language is English or All.
+		if ($lang !== 'en' && $lang != '*')
+		{
+			return $token;
+		}
+
+		// Stem the token if it is not in the cache.
+		if (!isset($this->cache[$lang][$token]))
+		{
+			// Stem the token.
+			$result = $token;
+			$result = self::_step1ab($result);
+			$result = self::_step1c($result);
+			$result = self::_step2($result);
+			$result = self::_step3($result);
+			$result = self::_step4($result);
+			$result = self::_step5($result);
+
+			// Add the token to the cache.
+			$this->cache[$lang][$token] = $result;
+		}
+
+		return $this->cache[$lang][$token];
+	}
+
+	/**
+	 * Step 1
+	 *
+	 * @param   string  $word  The token to stem.
+	 *
+	 * @return  string
+	 *
+	 * @since   2.5
+	 */
+	private static function _step1ab($word)
+	{
+		// Part a
+		if (substr($word, -1) == 's')
+		{
+			self::_replace($word, 'sses', 'ss')
+			or self::_replace($word, 'ies', 'i')
+			or self::_replace($word, 'ss', 'ss')
+			or self::_replace($word, 's', '');
+		}
+
+		// Part b
+		if (substr($word, -2, 1) != 'e' or !self::_replace($word, 'eed', 'ee', 0))
+		{
+			// First rule
+			$v = self::$_regex_vowel;
+
+			// Words ending with ing and ed
+			// Note use of && and OR, for precedence reasons
+			if (preg_match("#$v+#", substr($word, 0, -3)) && self::_replace($word, 'ing', '')
+				or preg_match("#$v+#", substr($word, 0, -2)) && self::_replace($word, 'ed', ''))
+			{
+				// If one of above two test successful
+				if (!self::_replace($word, 'at', 'ate') and !self::_replace($word, 'bl', 'ble') and !self::_replace($word, 'iz', 'ize'))
+				{
+					// Double consonant ending
+					if (self::_doubleConsonant($word) and substr($word, -2) != 'll' and substr($word, -2) != 'ss' and substr($word, -2) != 'zz')
+					{
+						$word = substr($word, 0, -1);
+					}
+					elseif (self::_m($word) == 1 and self::_cvc($word))
+					{
+						$word .= 'e';
+					}
+				}
+			}
+		}
+
+		return $word;
+	}
+
+	/**
+	 * Step 1c
+	 *
+	 * @param   string  $word  The token to stem.
+	 *
+	 * @return  string
+	 *
+	 * @since   2.5
+	 */
+	private static function _step1c($word)
+	{
+		$v = self::$_regex_vowel;
+
+		if (substr($word, -1) == 'y' && preg_match("#$v+#", substr($word, 0, -1)))
+		{
+			self::_replace($word, 'y', 'i');
+		}
+
+		return $word;
+	}
+
+	/**
+	 * Step 2
+	 *
+	 * @param   string  $word  The token to stem.
+	 *
+	 * @return  string
+	 *
+	 * @since   2.5
+	 */
+	private static function _step2($word)
+	{
+		switch (substr($word, -2, 1))
+		{
+			case 'a':
+				self::_replace($word, 'ational', 'ate', 0)
+				or self::_replace($word, 'tional', 'tion', 0);
+				break;
+			case 'c':
+				self::_replace($word, 'enci', 'ence', 0)
+				or self::_replace($word, 'anci', 'ance', 0);
+				break;
+			case 'e':
+				self::_replace($word, 'izer', 'ize', 0);
+				break;
+			case 'g':
+				self::_replace($word, 'logi', 'log', 0);
+				break;
+			case 'l':
+				self::_replace($word, 'entli', 'ent', 0)
+				or self::_replace($word, 'ousli', 'ous', 0)
+				or self::_replace($word, 'alli', 'al', 0)
+				or self::_replace($word, 'bli', 'ble', 0)
+				or self::_replace($word, 'eli', 'e', 0);
+				break;
+			case 'o':
+				self::_replace($word, 'ization', 'ize', 0)
+				or self::_replace($word, 'ation', 'ate', 0)
+				or self::_replace($word, 'ator', 'ate', 0);
+				break;
+			case 's':
+				self::_replace($word, 'iveness', 'ive', 0)
+				or self::_replace($word, 'fulness', 'ful', 0)
+				or self::_replace($word, 'ousness', 'ous', 0)
+				or self::_replace($word, 'alism', 'al', 0);
+				break;
+			case 't':
+				self::_replace($word, 'biliti', 'ble', 0)
+				or self::_replace($word, 'aliti', 'al', 0)
+				or self::_replace($word, 'iviti', 'ive', 0);
+				break;
+		}
+
+		return $word;
+	}
+
+	/**
+	 * Step 3
+	 *
+	 * @param   string  $word  The token to stem.
+	 *
+	 * @return  string
+	 *
+	 * @since   2.5
+	 */
+	private static function _step3($word)
+	{
+		switch (substr($word, -2, 1))
+		{
+			case 'a':
+				self::_replace($word, 'ical', 'ic', 0);
+				break;
+			case 's':
+				self::_replace($word, 'ness', '', 0);
+				break;
+			case 't':
+				self::_replace($word, 'icate', 'ic', 0)
+				or self::_replace($word, 'iciti', 'ic', 0);
+				break;
+			case 'u':
+				self::_replace($word, 'ful', '', 0);
+				break;
+			case 'v':
+				self::_replace($word, 'ative', '', 0);
+				break;
+			case 'z':
+				self::_replace($word, 'alize', 'al', 0);
+				break;
+		}
+
+		return $word;
+	}
+
+	/**
+	 * Step 4
+	 *
+	 * @param   string  $word  The token to stem.
+	 *
+	 * @return  string
+	 *
+	 * @since   2.5
+	 */
+	private static function _step4($word)
+	{
+		switch (substr($word, -2, 1))
+		{
+			case 'a':
+				self::_replace($word, 'al', '', 1);
+				break;
+			case 'c':
+					self::_replace($word, 'ance', '', 1)
+				or self::_replace($word, 'ence', '', 1);
+				break;
+			case 'e':
+				self::_replace($word, 'er', '', 1);
+				break;
+			case 'i':
+				self::_replace($word, 'ic', '', 1);
+				break;
+			case 'l':
+				self::_replace($word, 'able', '', 1)
+				or self::_replace($word, 'ible', '', 1);
+				break;
+			case 'n':
+				self::_replace($word, 'ant', '', 1)
+				or self::_replace($word, 'ement', '', 1)
+				or self::_replace($word, 'ment', '', 1)
+				or self::_replace($word, 'ent', '', 1);
+				break;
+			case 'o':
+				if (substr($word, -4) == 'tion' or substr($word, -4) == 'sion')
+				{
+					self::_replace($word, 'ion', '', 1);
+				}
+				else
+				{
+					self::_replace($word, 'ou', '', 1);
+				}
+				break;
+			case 's':
+				self::_replace($word, 'ism', '', 1);
+				break;
+			case 't':
+					self::_replace($word, 'ate', '', 1)
+				or self::_replace($word, 'iti', '', 1);
+				break;
+			case 'u':
+				self::_replace($word, 'ous', '', 1);
+				break;
+			case 'v':
+				self::_replace($word, 'ive', '', 1);
+				break;
+			case 'z':
+				self::_replace($word, 'ize', '', 1);
+				break;
+		}
+
+		return $word;
+	}
+
+	/**
+	 * Step 5
+	 *
+	 * @param   string  $word  The token to stem.
+	 *
+	 * @return  string
+	 *
+	 * @since   2.5
+	 */
+	private static function _step5($word)
+	{
+		// Part a
+		if (substr($word, -1) == 'e')
+		{
+			if (self::_m(substr($word, 0, -1)) > 1)
+			{
+				self::_replace($word, 'e', '');
+			}
+			elseif (self::_m(substr($word, 0, -1)) == 1)
+			{
+				if (!self::_cvc(substr($word, 0, -1)))
+				{
+					self::_replace($word, 'e', '');
+				}
+			}
+		}
+
+		// Part b
+		if (self::_m($word) > 1 and self::_doubleConsonant($word) and substr($word, -1) == 'l')
+		{
+			$word = substr($word, 0, -1);
+		}
+
+		return $word;
+	}
+
+	/**
+	 * Replaces the first string with the second, at the end of the string. If third
+	 * arg is given, then the preceding string must match that m count at least.
+	 *
+	 * @param   string   &$str   String to check
+	 * @param   string   $check  Ending to check for
+	 * @param   string   $repl   Replacement string
+	 * @param   integer  $m      Optional minimum number of m() to meet
+	 *
+	 * @return  boolean  Whether the $check string was at the end
+	 *                   of the $str string. True does not necessarily mean
+	 *                   that it was replaced.
+	 *
+	 * @since   2.5
+	 */
+	private static function _replace(&$str, $check, $repl, $m = null)
+	{
+		$len = 0 - strlen($check);
+
+		if (substr($str, $len) == $check)
+		{
+			$substr = substr($str, 0, $len);
+
+			if (is_null($m) or self::_m($substr) > $m)
+			{
+				$str = $substr . $repl;
+			}
+
+			return true;
+		}
+
+		return false;
+	}
+
+	/**
+	 * m() measures the number of consonant sequences in $str. if c is
+	 * a consonant sequence and v a vowel sequence, and <..> indicates arbitrary
+	 * presence,
+	 *
+	 * <c><v>       gives 0
+	 * <c>vc<v>     gives 1
+	 * <c>vcvc<v>   gives 2
+	 * <c>vcvcvc<v> gives 3
+	 *
+	 * @param   string  $str  The string to return the m count for
+	 *
+	 * @return  integer  The m count
+	 *
+	 * @since   2.5
+	 */
+	private static function _m($str)
+	{
+		$c = self::$_regex_consonant;
+		$v = self::$_regex_vowel;
+
+		$str = preg_replace("#^$c+#", '', $str);
+		$str = preg_replace("#$v+$#", '', $str);
+
+		preg_match_all("#($v+$c+)#", $str, $matches);
+
+		return count($matches[1]);
+	}
+
+	/**
+	 * Returns true/false as to whether the given string contains two
+	 * of the same consonant next to each other at the end of the string.
+	 *
+	 * @param   string  $str  String to check
+	 *
+	 * @return  boolean  Result
+	 *
+	 * @since   2.5
+	 */
+	private static function _doubleConsonant($str)
+	{
+		$c = self::$_regex_consonant;
+
+		return preg_match("#$c{2}$#", $str, $matches) and $matches[0]{0} == $matches[0]{1};
+	}
+
+	/**
+	 * Checks for ending CVC sequence where second C is not W, X or Y
+	 *
+	 * @param   string  $str  String to check
+	 *
+	 * @return  boolean  Result
+	 *
+	 * @since   2.5
+	 */
+	private static function _cvc($str)
+	{
+		$c = self::$_regex_consonant;
+		$v = self::$_regex_vowel;
+
+		return preg_match("#($c$v$c)$#", $str, $matches) and strlen($matches[1]) == 3 and $matches[1]{2} != 'w' and $matches[1]{2} != 'x'
+			and $matches[1]{2} != 'y';
+	}
+}
--- a/administrator/components/com_finder/helpers/indexer/stemmer/snowball.php
+++ b/administrator/components/com_finder/helpers/indexer/stemmer/snowball.php
@@ -0,0 +1,135 @@
+<?php
+/**
+ * @package     Joomla.Administrator
+ * @subpackage  com_finder
+ *
+ * @copyright   Copyright (C) 2005 - 2013 Open Source Matters, Inc. All rights reserved.
+ * @license     GNU General Public License version 2 or later; see LICENSE
+ */
+
+defined('_JEXEC') or die;
+
+JLoader::register('FinderIndexerStemmer', dirname(__DIR__) . '/stemmer.php');
+
+/**
+ * Snowball stemmer class for the Finder indexer package.
+ *
+ * @package     Joomla.Administrator
+ * @subpackage  com_finder
+ * @since       2.5
+ */
+class FinderIndexerStemmerSnowball extends FinderIndexerStemmer
+{
+	/**
+	 * Method to stem a token and return the root.
+	 *
+	 * @param   string  $token  The token to stem.
+	 * @param   string  $lang   The language of the token.
+	 *
+	 * @return  string  The root token.
+	 *
+	 * @since   2.5
+	 */
+	public function stem($token, $lang)
+	{
+		// Language to use if All is specified.
+		static $defaultLang = '';
+
+		// If language is All then try to get site default language.
+		if ($lang == '*' && $defaultLang == '')
+		{
+			$languages = JLanguageHelper::getLanguages();
+			$defaultLang = isset($languages[0]->sef) ? $languages[0]->sef : '*';
+			$lang = $defaultLang;
+		}
+
+		// Stem the token if it is not in the cache.
+		if (!isset($this->cache[$lang][$token]))
+		{
+			// Get the stem function from the language string.
+			switch ($lang)
+			{
+				// Danish stemmer.
+				case 'da':
+					$function = 'stem_danish';
+					break;
+
+				// German stemmer.
+				case 'de':
+					$function = 'stem_german';
+					break;
+
+				// English stemmer.
+				default:
+				case 'en':
+					$function = 'stem_english';
+					break;
+
+				// Spanish stemmer.
+				case 'es':
+					$function = 'stem_spanish';
+					break;
+
+				// Finnish stemmer.
+				case 'fi':
+					$function = 'stem_finnish';
+					break;
+
+				// French stemmer.
+				case 'fr':
+					$function = 'stem_french';
+					break;
+
+				// Hungarian stemmer.
+				case 'hu':
+					$function = 'stem_hungarian';
+					break;
+
+				// Italian stemmer.
+				case 'it':
+					$function = 'stem_italian';
+					break;
+
+				// Norwegian stemmer.
+				case 'nb':
+					$function = 'stem_norwegian';
+					break;
+
+				// Dutch stemmer.
+				case 'nl':
+					$function = 'stem_dutch';
+					break;
+
+				// Portuguese stemmer.
+				case 'pt':
+					$function = 'stem_portuguese';
+					break;
+
+				// Romanian stemmer.
+				case 'ro':
+					$function = 'stem_romanian';
+					break;
+
+				// Russian stemmer.
+				case 'ru':
+					$function = 'stem_russian_unicode';
+					break;
+
+				// Swedish stemmer.
+				case 'sv':
+					$function = 'stem_swedish';
+					break;
+
+				// Turkish stemmer.
+				case 'tr':
+					$function = 'stem_turkish_unicode';
+					break;
+			}
+
+			// Stem the word if the stemmer method exists.
+			$this->cache[$lang][$token] = function_exists($function) ? $function($token) : $token;
+		}
+
+		return $this->cache[$lang][$token];
+	}
+}