Current Path: > > opt > > alt > python39 > lib > python3.9 > site-packages > pip > _vendor > chardet
Operation : Linux premium131.web-hosting.com 4.18.0-553.44.1.lve.el8.x86_64 #1 SMP Thu Mar 13 14:29:12 UTC 2025 x86_64 Software : Apache Server IP : 162.0.232.56 | Your IP: 216.73.216.111 Domains : 1034 Domain(s) Permission : [ 0755 ]
Name | Type | Size | Last Modified | Actions |
---|---|---|---|---|
__pycache__ | Directory | - | - | |
cli | Directory | - | - | |
metadata | Directory | - | - | |
__init__.py | File | 3271 bytes | November 13 2023 21:47:44. | |
big5freq.py | File | 31254 bytes | November 13 2023 21:47:44. | |
big5prober.py | File | 1757 bytes | November 13 2023 21:47:44. | |
chardistribution.py | File | 9411 bytes | November 13 2023 21:47:44. | |
charsetgroupprober.py | File | 3839 bytes | November 13 2023 21:47:44. | |
charsetprober.py | File | 5110 bytes | November 13 2023 21:47:44. | |
codingstatemachine.py | File | 3590 bytes | November 13 2023 21:47:44. | |
compat.py | File | 1200 bytes | November 13 2023 21:47:44. | |
cp949prober.py | File | 1855 bytes | November 13 2023 21:47:44. | |
enums.py | File | 1661 bytes | November 13 2023 21:47:44. | |
escprober.py | File | 3950 bytes | November 13 2023 21:47:44. | |
escsm.py | File | 10510 bytes | November 13 2023 21:47:44. | |
eucjpprober.py | File | 3749 bytes | November 13 2023 21:47:44. | |
euckrfreq.py | File | 13546 bytes | November 13 2023 21:47:44. | |
euckrprober.py | File | 1748 bytes | November 13 2023 21:47:44. | |
euctwfreq.py | File | 31621 bytes | November 13 2023 21:47:44. | |
euctwprober.py | File | 1747 bytes | November 13 2023 21:47:44. | |
gb2312freq.py | File | 20715 bytes | November 13 2023 21:47:44. | |
gb2312prober.py | File | 1754 bytes | November 13 2023 21:47:44. | |
hebrewprober.py | File | 13838 bytes | November 13 2023 21:47:44. | |
jisfreq.py | File | 25777 bytes | November 13 2023 21:47:44. | |
jpcntx.py | File | 19643 bytes | November 13 2023 21:47:44. | |
langbulgarianmodel.py | File | 105675 bytes | November 13 2023 21:47:44. | |
langgreekmodel.py | File | 99549 bytes | November 13 2023 21:47:44. | |
langhebrewmodel.py | File | 98754 bytes | November 13 2023 21:47:44. | |
langhungarianmodel.py | File | 102476 bytes | November 13 2023 21:47:44. | |
langrussianmodel.py | File | 131158 bytes | November 13 2023 21:47:44. | |
langthaimodel.py | File | 103290 bytes | November 13 2023 21:47:44. | |
langturkishmodel.py | File | 95924 bytes | November 13 2023 21:47:44. | |
latin1prober.py | File | 5370 bytes | November 13 2023 21:47:44. | |
mbcharsetprober.py | File | 3413 bytes | November 13 2023 21:47:44. | |
mbcsgroupprober.py | File | 2012 bytes | November 13 2023 21:47:44. | |
mbcssm.py | File | 25481 bytes | November 13 2023 21:47:44. | |
sbcharsetprober.py | File | 6136 bytes | November 13 2023 21:47:44. | |
sbcsgroupprober.py | File | 4309 bytes | November 13 2023 21:47:44. | |
sjisprober.py | File | 3774 bytes | November 13 2023 21:47:44. | |
universaldetector.py | File | 12503 bytes | November 13 2023 21:47:44. | |
utf8prober.py | File | 2766 bytes | November 13 2023 21:47:44. | |
version.py | File | 242 bytes | November 13 2023 21:47:44. |
######################## BEGIN LICENSE BLOCK ######################## # The Original Code is Mozilla Universal charset detector code. # # The Initial Developer of the Original Code is # Netscape Communications Corporation. # Portions created by the Initial Developer are Copyright (C) 2001 # the Initial Developer. All Rights Reserved. # # Contributor(s): # Mark Pilgrim - port to Python # Shy Shalom - original C code # # This library is free software; you can redistribute it and/or # modify it under the terms of the GNU Lesser General Public # License as published by the Free Software Foundation; either # version 2.1 of the License, or (at your option) any later version. # # This library is distributed in the hope that it will be useful, # but WITHOUT ANY WARRANTY; without even the implied warranty of # MERCHANTABILITY or FITNESS FOR A PARTICULAR PURPOSE. See the GNU # Lesser General Public License for more details. # # You should have received a copy of the GNU Lesser General Public # License along with this library; if not, write to the Free Software # Foundation, Inc., 51 Franklin St, Fifth Floor, Boston, MA # 02110-1301 USA ######################### END LICENSE BLOCK ######################### from .charsetgroupprober import CharSetGroupProber from .hebrewprober import HebrewProber from .langbulgarianmodel import (ISO_8859_5_BULGARIAN_MODEL, WINDOWS_1251_BULGARIAN_MODEL) from .langgreekmodel import ISO_8859_7_GREEK_MODEL, WINDOWS_1253_GREEK_MODEL from .langhebrewmodel import WINDOWS_1255_HEBREW_MODEL # from .langhungarianmodel import (ISO_8859_2_HUNGARIAN_MODEL, # WINDOWS_1250_HUNGARIAN_MODEL) from .langrussianmodel import (IBM855_RUSSIAN_MODEL, IBM866_RUSSIAN_MODEL, ISO_8859_5_RUSSIAN_MODEL, KOI8_R_RUSSIAN_MODEL, MACCYRILLIC_RUSSIAN_MODEL, WINDOWS_1251_RUSSIAN_MODEL) from .langthaimodel import TIS_620_THAI_MODEL from .langturkishmodel import ISO_8859_9_TURKISH_MODEL from .sbcharsetprober import SingleByteCharSetProber class SBCSGroupProber(CharSetGroupProber): def __init__(self): super(SBCSGroupProber, self).__init__() hebrew_prober = HebrewProber() logical_hebrew_prober = SingleByteCharSetProber(WINDOWS_1255_HEBREW_MODEL, False, hebrew_prober) # TODO: See if using ISO-8859-8 Hebrew model works better here, since # it's actually the visual one visual_hebrew_prober = SingleByteCharSetProber(WINDOWS_1255_HEBREW_MODEL, True, hebrew_prober) hebrew_prober.set_model_probers(logical_hebrew_prober, visual_hebrew_prober) # TODO: ORDER MATTERS HERE. I changed the order vs what was in master # and several tests failed that did not before. Some thought # should be put into the ordering, and we should consider making # order not matter here, because that is very counter-intuitive. self.probers = [ SingleByteCharSetProber(WINDOWS_1251_RUSSIAN_MODEL), SingleByteCharSetProber(KOI8_R_RUSSIAN_MODEL), SingleByteCharSetProber(ISO_8859_5_RUSSIAN_MODEL), SingleByteCharSetProber(MACCYRILLIC_RUSSIAN_MODEL), SingleByteCharSetProber(IBM866_RUSSIAN_MODEL), SingleByteCharSetProber(IBM855_RUSSIAN_MODEL), SingleByteCharSetProber(ISO_8859_7_GREEK_MODEL), SingleByteCharSetProber(WINDOWS_1253_GREEK_MODEL), SingleByteCharSetProber(ISO_8859_5_BULGARIAN_MODEL), SingleByteCharSetProber(WINDOWS_1251_BULGARIAN_MODEL), # TODO: Restore Hungarian encodings (iso-8859-2 and windows-1250) # after we retrain model. # SingleByteCharSetProber(ISO_8859_2_HUNGARIAN_MODEL), # SingleByteCharSetProber(WINDOWS_1250_HUNGARIAN_MODEL), SingleByteCharSetProber(TIS_620_THAI_MODEL), SingleByteCharSetProber(ISO_8859_9_TURKISH_MODEL), hebrew_prober, logical_hebrew_prober, visual_hebrew_prober, ] self.reset()
SILENT KILLER Tool