dosage/dosagelib/util.py

# -*- coding: iso-8859-1 -*-
# Copyright (C) 2004-2005 Tristan Seligmann and Jonathan Jacobs
# Copyright (C) 2012 Bastian Kleineidam
from __future__ import division, print_function

import urllib, urllib2, urlparse
import requests
import sys
import os
import cgi
import re
import traceback
import time
from htmlentitydefs import name2codepoint

from .output import out
from .configuration import UserAgent, AppName, App, SupportUrl
from .fileutil import has_module, is_tty
if os.name == 'nt':
    from . import colorama

has_curses = has_module("curses")

MAX_FILESIZE = 1024*1024*1 # 1MB

def tagre(tag, attribute, value, quote='"', before="", after=""):
    """Return a regular expression matching the given HTML tag, attribute
    and value. It matches the tag and attribute names case insensitive,
    and skips arbitrary whitespace and leading HTML attributes. The "<>" at
    the start and end of the HTML tag is also matched.
    @param tag: the tag name
    @ptype tag: string
    @param attribute: the attribute name
    @ptype attribute: string
    @param value: the attribute value
    @ptype value: string
    @param quote: the attribute quote (default ")
    @ptype quote: string
    @param after: match after attribute value but before end
    @ptype after: string

    @return: the generated regular expression suitable for re.compile()
    @rtype: string
    """
    if before:
        prefix = r"[^>]*%s[^>]*\s+" % before
    else:
        prefix = r"(?:[^>]*\s+)?"
    attrs = dict(
        tag=case_insensitive_re(tag),
        attribute=case_insensitive_re(attribute),
        value=value,
        quote=quote,
        prefix=prefix,
        after=after,
    )
    return r'<\s*%(tag)s\s+%(prefix)s%(attribute)s\s*=\s*%(quote)s%(value)s%(quote)s[^>]*%(after)s[^>]*>' % attrs


def case_insensitive_re(name):
    """Reformat the given name to a case insensitive regular expression string
    without using re.IGNORECASE. This way selective strings can be made case
    insensitive.
    @param name: the name to make case insensitive
    @ptype name: string
    @return: the case insenstive regex
    @rtype: string
    """
    return "".join("[%s%s]" % (c.lower(), c.upper()) for c in name)


baseSearch = re.compile(tagre("base", "href", '([^"]*)'))

def getPageContent(url):
    # read page data
    page = urlopen(url)
    data = page.text
    # determine base URL
    baseUrl = None
    match = baseSearch.search(data)
    if match:
        baseUrl = match.group(1)
    else:
        baseUrl = url
    return data, baseUrl


def fetchUrl(url, urlSearch):
    data, baseUrl = getPageContent(url)
    match = urlSearch.search(data)
    if match:
        searchUrl = match.group(1)
        if not searchUrl:
            raise ValueError("Match empty URL at %s with pattern %s" % (url, urlSearch.pattern))
        out.write('matched URL %r' % searchUrl, 2)
        return normaliseURL(urlparse.urljoin(baseUrl, searchUrl))
    return None


def fetchUrls(url, imageSearch, prevSearch=None):
    data, baseUrl = getPageContent(url)
    # match images
    imageUrls = set()
    for match in imageSearch.finditer(data):
        imageUrl = match.group(1)
        if not imageUrl:
            raise ValueError("Match empty image URL at %s with pattern %s" % (url, imageSearch.pattern))
        out.write('matched image URL %r with pattern %s' % (imageUrl, imageSearch.pattern), 2)
        imageUrls.add(normaliseURL(urlparse.urljoin(baseUrl, imageUrl)))
    if not imageUrls:
        out.write("warning: no images found at %s with pattern %s" % (url, imageSearch.pattern))
    if prevSearch is not None:
        # match previous URL
        match = prevSearch.search(data)
        if match:
            prevUrl = match.group(1)
            if not prevUrl:
                raise ValueError("Match empty previous URL at %s with pattern %s" % (url, prevSearch.pattern))
            out.write('matched previous URL %r' % prevUrl, 2)
            prevUrl = normaliseURL(urlparse.urljoin(baseUrl, prevUrl))
        else:
            out.write('no previous URL %s at %s' % (prevSearch.pattern, url), 2)
            prevUrl = None
        return imageUrls, prevUrl
    return imageUrls, None


def unescape(text):
    """
    Replace HTML entities and character references.
    """
    def _fixup(m):
        text = m.group(0)
        if text[:2] == "&#":
            # character reference
            try:
                if text[:3] == "&#x":
                    text = unichr(int(text[3:-1], 16))
                else:
                    text = unichr(int(text[2:-1]))
            except ValueError:
                pass
        else:
            # named entity
            try:
                text = unichr(name2codepoint[text[1:-1]])
            except KeyError:
                pass
        if isinstance(text, unicode):
            text = text.encode('utf-8')
            text = urllib2.quote(text, safe=';/?:@&=+$,')
        return text
    return re.sub(r"&#?\w+;", _fixup, text)


def normaliseURL(url):
    """
    Removes any leading empty segments to avoid breaking urllib2; also replaces
    HTML entities and character references.
    """
    # XXX: brutal hack
    url = unescape(url)

    pu = list(urlparse.urlparse(url))
    segments = pu[2].split('/')
    while segments and segments[0] == '':
        del segments[0]
    pu[2] = '/' + '/'.join(segments).replace(' ', '%20')
    # remove leading '&' from query
    if pu[4].startswith('&'):
        pu[4] = pu[4][1:]
    # remove anchor
    pu[5] = ""
    return urlparse.urlunparse(pu)


def urlopen(url, referrer=None, retries=3, retry_wait_seconds=5):
    out.write('Open URL %s' % url, 2)
    assert retries >= 0, 'invalid retry value %r' % retries
    assert retry_wait_seconds > 0, 'invalid retry seconds value %r' % retry_wait_seconds
    headers = {'User-Agent': UserAgent}
    config = {"max_retries": retries}
    if referrer:
        headers['Referer'] = referrer
    try:
        req = requests.get(url, headers=headers, config=config)
        req.raise_for_status()
        return req
    except requests.exceptions.RequestException as err:
        msg = 'URL retrieval of %s failed: %s' % (url, err)
        out.write(msg)
        raise IOError(msg)


def get_columns (fp):
    """Return number of columns for given file."""
    if not is_tty(fp):
        return 80
    if os.name == 'nt':
        return colorama.get_console_size().X
    if has_curses:
        import curses
        try:
            curses.setupterm(os.environ.get("TERM"), fp.fileno())
            return curses.tigetnum("cols")
        except curses.error:
           pass
    return 80


def splitpath(path):
    c = []
    head, tail = os.path.split(path)
    while tail:
        c.insert(0, tail)
        head, tail = os.path.split(head)
    return c


def getRelativePath(basepath, path):
    basepath = splitpath(os.path.abspath(basepath))
    path = splitpath(os.path.abspath(path))
    afterCommon = False
    for c in basepath:
        if afterCommon or path[0] != c:
            path.insert(0, os.path.pardir)
            afterCommon = True
        else:
            del path[0]
    return os.path.join(*path)


def getQueryParams(url):
    query = urlparse.urlsplit(url)[3]
    out.write('Extracting query parameters from %r (%r)...' % (url, query), 3)
    return cgi.parse_qs(query)


def internal_error(out=sys.stderr, etype=None, evalue=None, tb=None):
    """Print internal error message (output defaults to stderr)."""
    print(os.linesep, file=out)
    print("""********** Oops, I did it again. *************

You have found an internal error in %(app)s. Please write a bug report
at %(url)s and include at least the information below:

Not disclosing some of the information below due to privacy reasons is ok.
I will try to help you nonetheless, but you have to give me something
I can work with ;) .
""" % dict(app=AppName, url=SupportUrl), file=out)
    if etype is None:
        etype = sys.exc_info()[0]
    if evalue is None:
        evalue = sys.exc_info()[1]
    print(etype, evalue, file=out)
    if tb is None:
        tb = sys.exc_info()[2]
    traceback.print_exception(etype, evalue, tb, None, out)
    print_app_info(out=out)
    print_proxy_info(out=out)
    print_locale_info(out=out)
    print(os.linesep,
            "******** %s internal error, over and out ********" % AppName, file=out)


def print_env_info(key, out=sys.stderr):
    """If given environment key is defined, print it out."""
    value = os.getenv(key)
    if value is not None:
        print(key, "=", repr(value), file=out)


def print_proxy_info(out=sys.stderr):
    """Print proxy info."""
    print_env_info("http_proxy", out=out)


def print_locale_info(out=sys.stderr):
    """Print locale info."""
    for key in ("LANGUAGE", "LC_ALL", "LC_CTYPE", "LANG"):
        print_env_info(key, out=out)


def print_app_info(out=sys.stderr):
    """Print system and application info (output defaults to stderr)."""
    print("System info:", file=out)
    print(App, file=out)
    print("Python %(version)s on %(platform)s" %
                    {"version": sys.version, "platform": sys.platform}, file=out)
    stime = strtime(time.time())
    print("Local time:", stime, file=out)
    print("sys.argv", sys.argv, file=out)


def strtime(t):
    """Return ISO 8601 formatted time."""
    return time.strftime("%Y-%m-%d %H:%M:%S", time.localtime(t)) + \
           strtimezone()


def strtimezone():
    """Return timezone info, %z on some platforms, but not supported on all.
    """
    if time.daylight:
        zone = time.altzone
    else:
        zone = time.timezone
    return "%+04d" % (-zone//3600)


def asciify(name):
    """Remove non-ascii characters from string."""
    return re.sub("[^0-9a-zA-Z_]", "", name)


def unquote(text):
    while '%' in text:
        text = urllib.unquote(text)
    return text


def strsize (b):
    """Return human representation of bytes b. A negative number of bytes
    raises a value error."""
    if b < 0:
        raise ValueError("Invalid negative byte number")
    if b < 1024:
        return "%dB" % b
    if b < 1024 * 10:
        return "%dKB" % (b // 1024)
    if b < 1024 * 1024:
        return "%.2fKB" % (float(b) / 1024)
    if b < 1024 * 1024 * 10:
        return "%.2fMB" % (float(b) / (1024*1024))
    if b < 1024 * 1024 * 1024:
        return "%.1fMB" % (float(b) / (1024*1024))
    if b < 1024 * 1024 * 1024 * 10:
        return "%.2fGB" % (float(b) / (1024*1024*1024))
    return "%.1fGB" % (float(b) / (1024*1024*1024))
Updated copyright for all source files. 2012-06-20 20:41:04 +00:00			`# -- coding: iso-8859-1 --`
			`# Copyright (C) 2004-2005 Tristan Seligmann and Jonathan Jacobs`
			`# Copyright (C) 2012 Bastian Kleineidam`
Updated documentation and fix some comics. 2012-11-20 17:53:53 +00:00			`from __future__ import division, print_function`
Initial commit to Github. 2012-06-20 19:58:13 +00:00
Add comic scripts, add fixes and other stuff. 2012-11-28 17:15:12 +00:00			`import urllib, urllib2, urlparse`
Fix comics, improve tests, use python-requests. 2012-11-26 17:44:31 +00:00			`import requests`
Initial commit to Github. 2012-06-20 19:58:13 +00:00			`import sys`
			`import os`
			`import cgi`
			`import re`
			`import traceback`
			`import time`
			`from htmlentitydefs import name2codepoint`

			`from .output import out`
			`from .configuration import UserAgent, AppName, App, SupportUrl`
Improved terminal functions. 2012-06-20 20:33:26 +00:00			`from .fileutil import has_module, is_tty`
Only import colorama on windows systems. 2012-10-01 16:01:56 +00:00			`if os.name == 'nt':`
			`from . import colorama`
Improved terminal functions. 2012-06-20 20:33:26 +00:00
			`has_curses = has_module("curses")`
Initial commit to Github. 2012-06-20 19:58:13 +00:00
A lot of refactoring. 2012-10-11 10:03:12 +00:00			`MAX_FILESIZE = 102410241 # 1MB`
Initial commit to Github. 2012-06-20 19:58:13 +00:00
Match before and after a tag. 2012-10-12 19:11:44 +00:00			`def tagre(tag, attribute, value, quote='"', before="", after=""):`
A lot of refactoring. 2012-10-11 10:03:12 +00:00			`"""Return a regular expression matching the given HTML tag, attribute`
			`and value. It matches the tag and attribute names case insensitive,`
			`and skips arbitrary whitespace and leading HTML attributes. The "<>" at`
			`the start and end of the HTML tag is also matched.`
			`@param tag: the tag name`
			`@ptype tag: string`
			`@param attribute: the attribute name`
			`@ptype attribute: string`
			`@param value: the attribute value`
			`@ptype value: string`
Make tagre quote configurable. 2012-10-11 13:43:29 +00:00			`@param quote: the attribute quote (default ")`
			`@ptype quote: string`
Match before and after a tag. 2012-10-12 19:11:44 +00:00			`@param after: match after attribute value but before end`
			`@ptype after: string`

A lot of refactoring. 2012-10-11 10:03:12 +00:00			`@return: the generated regular expression suitable for re.compile()`
			`@rtype: string`
			`"""`
Add comic scripts, add fixes and other stuff. 2012-11-28 17:15:12 +00:00			`if before:`
			`prefix = r"[^>]%s[^>]\s+" % before`
			`else:`
			`prefix = r"(?:[^>]*\s+)?"`
A lot of refactoring. 2012-10-11 10:03:12 +00:00			`attrs = dict(`
			`tag=case_insensitive_re(tag),`
			`attribute=case_insensitive_re(attribute),`
			`value=value,`
Make tagre quote configurable. 2012-10-11 13:43:29 +00:00			`quote=quote,`
Add comic scripts, add fixes and other stuff. 2012-11-28 17:15:12 +00:00			`prefix=prefix,`
Match before and after a tag. 2012-10-12 19:11:44 +00:00			`after=after,`
A lot of refactoring. 2012-10-11 10:03:12 +00:00			`)`
Add comic scripts, add fixes and other stuff. 2012-11-28 17:15:12 +00:00			`return r'<\s%(tag)s\s+%(prefix)s%(attribute)s\s=\s%(quote)s%(value)s%(quote)s[^>]%(after)s[^>]*>' % attrs`
Initial commit to Github. 2012-06-20 19:58:13 +00:00

A lot of refactoring. 2012-10-11 10:03:12 +00:00			`def case_insensitive_re(name):`
			`"""Reformat the given name to a case insensitive regular expression string`
			`without using re.IGNORECASE. This way selective strings can be made case`
			`insensitive.`
			`@param name: the name to make case insensitive`
			`@ptype name: string`
			`@return: the case insenstive regex`
			`@rtype: string`
			`"""`
			`return "".join("[%s%s]" % (c.lower(), c.upper()) for c in name)`
Initial commit to Github. 2012-06-20 19:58:13 +00:00

A lot of refactoring. 2012-10-11 10:03:12 +00:00			`baseSearch = re.compile(tagre("base", "href", '([^"]*)'))`
Initial commit to Github. 2012-06-20 19:58:13 +00:00
A lot of refactoring. 2012-10-11 10:03:12 +00:00			`def getPageContent(url):`
			`# read page data`
			`page = urlopen(url)`
Fix comics, improve tests, use python-requests. 2012-11-26 17:44:31 +00:00			`data = page.text`
A lot of refactoring. 2012-10-11 10:03:12 +00:00			`# determine base URL`
			`baseUrl = None`
			`match = baseSearch.search(data)`
			`if match:`
			`baseUrl = match.group(1)`
Initial commit to Github. 2012-06-20 19:58:13 +00:00			`else:`
A lot of refactoring. 2012-10-11 10:03:12 +00:00			`baseUrl = url`
			`return data, baseUrl`
Initial commit to Github. 2012-06-20 19:58:13 +00:00

Prevent empty URL matching. 2012-10-11 16:16:29 +00:00			`def fetchUrl(url, urlSearch):`
A lot of refactoring. 2012-10-11 10:03:12 +00:00			`data, baseUrl = getPageContent(url)`
Prevent empty URL matching. 2012-10-11 16:16:29 +00:00			`match = urlSearch.search(data)`
Initial commit to Github. 2012-06-20 19:58:13 +00:00			`if match:`
A lot of refactoring. 2012-10-11 10:03:12 +00:00			`searchUrl = match.group(1)`
Prevent empty URL matching. 2012-10-11 16:16:29 +00:00			`if not searchUrl:`
			`raise ValueError("Match empty URL at %s with pattern %s" % (url, urlSearch.pattern))`
A lot of refactoring. 2012-10-11 10:03:12 +00:00			`out.write('matched URL %r' % searchUrl, 2)`
Fix some comics. 2012-11-21 20:57:26 +00:00			`return normaliseURL(urlparse.urljoin(baseUrl, searchUrl))`
Initial commit to Github. 2012-06-20 19:58:13 +00:00			`return None`


A lot of refactoring. 2012-10-11 10:03:12 +00:00			`def fetchUrls(url, imageSearch, prevSearch=None):`
			`data, baseUrl = getPageContent(url)`
			`# match images`
			`imageUrls = set()`
			`for match in imageSearch.finditer(data):`
			`imageUrl = match.group(1)`
Prevent empty URL matching. 2012-10-11 16:16:29 +00:00			`if not imageUrl:`
			`raise ValueError("Match empty image URL at %s with pattern %s" % (url, imageSearch.pattern))`
Fix comics, improve tests, use python-requests. 2012-11-26 17:44:31 +00:00			`out.write('matched image URL %r with pattern %s' % (imageUrl, imageSearch.pattern), 2)`
Fix some comics. 2012-11-21 20:57:26 +00:00			`imageUrls.add(normaliseURL(urlparse.urljoin(baseUrl, imageUrl)))`
A lot of refactoring. 2012-10-11 10:03:12 +00:00			`if not imageUrls:`
Only warn about missing images. 2012-10-11 13:17:08 +00:00			`out.write("warning: no images found at %s with pattern %s" % (url, imageSearch.pattern))`
A lot of refactoring. 2012-10-11 10:03:12 +00:00			`if prevSearch is not None:`
			`# match previous URL`
			`match = prevSearch.search(data)`
			`if match:`
			`prevUrl = match.group(1)`
Prevent empty URL matching. 2012-10-11 16:16:29 +00:00			`if not prevUrl:`
			`raise ValueError("Match empty previous URL at %s with pattern %s" % (url, prevSearch.pattern))`
A lot of refactoring. 2012-10-11 10:03:12 +00:00			`out.write('matched previous URL %r' % prevUrl, 2)`
Fix some comics. 2012-11-21 20:57:26 +00:00			`prevUrl = normaliseURL(urlparse.urljoin(baseUrl, prevUrl))`
A lot of refactoring. 2012-10-11 10:03:12 +00:00			`else:`
Make tagre quote configurable. 2012-10-11 13:43:29 +00:00			`out.write('no previous URL %s at %s' % (prevSearch.pattern, url), 2)`
A lot of refactoring. 2012-10-11 10:03:12 +00:00			`prevUrl = None`
			`return imageUrls, prevUrl`
Fix some comics. 2012-11-21 20:57:26 +00:00			`return imageUrls, None`
Initial commit to Github. 2012-06-20 19:58:13 +00:00

Add comic scripts, add fixes and other stuff. 2012-11-28 17:15:12 +00:00			`def unescape(text):`
Initial commit to Github. 2012-06-20 19:58:13 +00:00			`"""`
			`Replace HTML entities and character references.`
			`"""`
			`def _fixup(m):`
			`text = m.group(0)`
			`if text[:2] == "&#":`
			`# character reference`
			`try:`
			`if text[:3] == "&#x":`
			`text = unichr(int(text[3:-1], 16))`
			`else:`
			`text = unichr(int(text[2:-1]))`
			`except ValueError:`
			`pass`
			`else:`
			`# named entity`
			`try:`
			`text = unichr(name2codepoint[text[1:-1]])`
			`except KeyError:`
			`pass`
			`if isinstance(text, unicode):`
			`text = text.encode('utf-8')`
			`text = urllib2.quote(text, safe=';/?:@&=+$,')`
			`return text`
Fix some comics. 2012-11-21 20:57:26 +00:00			`return re.sub(r"&#?\w+;", _fixup, text)`

Initial commit to Github. 2012-06-20 19:58:13 +00:00
			`def normaliseURL(url):`
			`"""`
			`Removes any leading empty segments to avoid breaking urllib2; also replaces`
			`HTML entities and character references.`
			`"""`
			`# XXX: brutal hack`
Add comic scripts, add fixes and other stuff. 2012-11-28 17:15:12 +00:00			`url = unescape(url)`
Initial commit to Github. 2012-06-20 19:58:13 +00:00
			`pu = list(urlparse.urlparse(url))`
Fix some comics. 2012-11-21 20:57:26 +00:00			`segments = pu[2].split('/')`
Initial commit to Github. 2012-06-20 19:58:13 +00:00			`while segments and segments[0] == '':`
			`del segments[0]`
Fix some comics. 2012-11-21 20:57:26 +00:00			`pu[2] = '/' + '/'.join(segments).replace(' ', '%20')`
Fix some comics 2012-11-14 19:23:30 +00:00			`# remove leading '&' from query`
Fix some comics. 2012-11-21 20:57:26 +00:00			`if pu[4].startswith('&'):`
			`pu[4] = pu[4][1:]`
			`# remove anchor`
			`pu[5] = ""`
Initial commit to Github. 2012-06-20 19:58:13 +00:00			`return urlparse.urlunparse(pu)`

Fix some comics. 2012-11-21 20:57:26 +00:00
Improve URL retrieval. 2012-10-11 17:53:10 +00:00			`def urlopen(url, referrer=None, retries=3, retry_wait_seconds=5):`
			`out.write('Open URL %s' % url, 2)`
Code cleanup. 2012-09-27 19:24:28 +00:00			`assert retries >= 0, 'invalid retry value %r' % retries`
			`assert retry_wait_seconds > 0, 'invalid retry seconds value %r' % retry_wait_seconds`
Fix comics, improve tests, use python-requests. 2012-11-26 17:44:31 +00:00			`headers = {'User-Agent': UserAgent}`
			`config = {"max_retries": retries}`
Initial commit to Github. 2012-06-20 19:58:13 +00:00			`if referrer:`
Fix comics, improve tests, use python-requests. 2012-11-26 17:44:31 +00:00			`headers['Referer'] = referrer`
			`try:`
			`req = requests.get(url, headers=headers, config=config)`
			`req.raise_for_status()`
			`return req`
			`except requests.exceptions.RequestException as err:`
			`msg = 'URL retrieval of %s failed: %s' % (url, err)`
			`out.write(msg)`
			`raise IOError(msg)`
Initial commit to Github. 2012-06-20 19:58:13 +00:00
Improved terminal functions. 2012-06-20 20:33:26 +00:00
			`def get_columns (fp):`
			`"""Return number of columns for given file."""`
			`if not is_tty(fp):`
			`return 80`
Improve console size guessing. 2012-09-27 19:59:11 +00:00			`if os.name == 'nt':`
			`return colorama.get_console_size().X`
Improved terminal functions. 2012-06-20 20:33:26 +00:00			`if has_curses:`
			`import curses`
			`try:`
Improve console size guessing. 2012-09-27 19:59:11 +00:00			`curses.setupterm(os.environ.get("TERM"), fp.fileno())`
Improved terminal functions. 2012-06-20 20:33:26 +00:00			`return curses.tigetnum("cols")`
			`except curses.error:`
			`pass`
			`return 80`

Initial commit to Github. 2012-06-20 19:58:13 +00:00
			`def splitpath(path):`
			`c = []`
			`head, tail = os.path.split(path)`
			`while tail:`
			`c.insert(0, tail)`
			`head, tail = os.path.split(head)`
			`return c`

Fix some comics. 2012-11-21 20:57:26 +00:00
Initial commit to Github. 2012-06-20 19:58:13 +00:00			`def getRelativePath(basepath, path):`
			`basepath = splitpath(os.path.abspath(basepath))`
			`path = splitpath(os.path.abspath(path))`
			`afterCommon = False`
			`for c in basepath:`
			`if afterCommon or path[0] != c:`
			`path.insert(0, os.path.pardir)`
			`afterCommon = True`
			`else:`
			`del path[0]`
			`return os.path.join(*path)`

Fix some comics. 2012-11-21 20:57:26 +00:00
Initial commit to Github. 2012-06-20 19:58:13 +00:00			`def getQueryParams(url):`
			`query = urlparse.urlsplit(url)[3]`
			`out.write('Extracting query parameters from %r (%r)...' % (url, query), 3)`
			`return cgi.parse_qs(query)`


			`def internal_error(out=sys.stderr, etype=None, evalue=None, tb=None):`
			`"""Print internal error message (output defaults to stderr)."""`
Updated documentation and fix some comics. 2012-11-20 17:53:53 +00:00			`print(os.linesep, file=out)`
			`print("""******** Oops, I did it again. ***********`
Initial commit to Github. 2012-06-20 19:58:13 +00:00
			`You have found an internal error in %(app)s. Please write a bug report`
Fix comics, improve tests, use python-requests. 2012-11-26 17:44:31 +00:00			`at %(url)s and include at least the information below:`
Initial commit to Github. 2012-06-20 19:58:13 +00:00
Fix comics, improve tests, use python-requests. 2012-11-26 17:44:31 +00:00			`Not disclosing some of the information below due to privacy reasons is ok.`
Initial commit to Github. 2012-06-20 19:58:13 +00:00			`I will try to help you nonetheless, but you have to give me something`
			`I can work with ;) .`
Updated documentation and fix some comics. 2012-11-20 17:53:53 +00:00			`""" % dict(app=AppName, url=SupportUrl), file=out)`
Initial commit to Github. 2012-06-20 19:58:13 +00:00			`if etype is None:`
			`etype = sys.exc_info()[0]`
			`if evalue is None:`
			`evalue = sys.exc_info()[1]`
Fix some comics. 2012-11-21 20:57:26 +00:00			`print(etype, evalue, file=out)`
Initial commit to Github. 2012-06-20 19:58:13 +00:00			`if tb is None:`
			`tb = sys.exc_info()[2]`
			`traceback.print_exception(etype, evalue, tb, None, out)`
			`print_app_info(out=out)`
			`print_proxy_info(out=out)`
			`print_locale_info(out=out)`
Updated documentation and fix some comics. 2012-11-20 17:53:53 +00:00			`print(os.linesep,`
			`"****** %s internal error, over and out ******" % AppName, file=out)`
Initial commit to Github. 2012-06-20 19:58:13 +00:00

			`def print_env_info(key, out=sys.stderr):`
			`"""If given environment key is defined, print it out."""`
			`value = os.getenv(key)`
			`if value is not None:`
Updated documentation and fix some comics. 2012-11-20 17:53:53 +00:00			`print(key, "=", repr(value), file=out)`
Initial commit to Github. 2012-06-20 19:58:13 +00:00

			`def print_proxy_info(out=sys.stderr):`
			`"""Print proxy info."""`
			`print_env_info("http_proxy", out=out)`


			`def print_locale_info(out=sys.stderr):`
			`"""Print locale info."""`
			`for key in ("LANGUAGE", "LC_ALL", "LC_CTYPE", "LANG"):`
			`print_env_info(key, out=out)`


			`def print_app_info(out=sys.stderr):`
			`"""Print system and application info (output defaults to stderr)."""`
Updated documentation and fix some comics. 2012-11-20 17:53:53 +00:00			`print("System info:", file=out)`
			`print(App, file=out)`
			`print("Python %(version)s on %(platform)s" %`
			`{"version": sys.version, "platform": sys.platform}, file=out)`
Initial commit to Github. 2012-06-20 19:58:13 +00:00			`stime = strtime(time.time())`
Updated documentation and fix some comics. 2012-11-20 17:53:53 +00:00			`print("Local time:", stime, file=out)`
Fix comics, improve tests, use python-requests. 2012-11-26 17:44:31 +00:00			`print("sys.argv", sys.argv, file=out)`
Initial commit to Github. 2012-06-20 19:58:13 +00:00

			`def strtime(t):`
			`"""Return ISO 8601 formatted time."""`
			`return time.strftime("%Y-%m-%d %H:%M:%S", time.localtime(t)) + \`
			`strtimezone()`


			`def strtimezone():`
			`"""Return timezone info, %z on some platforms, but not supported on all.`
			`"""`
			`if time.daylight:`
			`zone = time.altzone`
			`else:`
			`zone = time.timezone`
			`return "%+04d" % (-zone//3600)`
Dynamic type generation helpers. 2012-11-26 06:14:02 +00:00

			`def asciify(name):`
			`"""Remove non-ascii characters from string."""`
			`return re.sub("[^0-9a-zA-Z_]", "", name)`
Add comic scripts, add fixes and other stuff. 2012-11-28 17:15:12 +00:00

			`def unquote(text):`
			`while '%' in text:`
			`text = urllib.unquote(text)`
			`return text`
Fix some comics. 2012-12-02 17:35:06 +00:00

			`def strsize (b):`
			`"""Return human representation of bytes b. A negative number of bytes`
			`raises a value error."""`
			`if b < 0:`
			`raise ValueError("Invalid negative byte number")`
			`if b < 1024:`
			`return "%dB" % b`
			`if b < 1024 * 10:`
			`return "%dKB" % (b // 1024)`
			`if b < 1024 * 1024:`
			`return "%.2fKB" % (float(b) / 1024)`
			`if b < 1024 * 1024 * 10:`
			`return "%.2fMB" % (float(b) / (1024*1024))`
			`if b < 1024 * 1024 * 1024:`
			`return "%.1fMB" % (float(b) / (1024*1024))`
			`if b < 1024 * 1024 * 1024 * 10:`
			`return "%.2fGB" % (float(b) / (102410241024))`
			`return "%.1fGB" % (float(b) / (102410241024))`