dosage/dosagelib/plugins/s.py

# -*- coding: iso-8859-1 -*-
# Copyright (C) 2004-2005 Tristan Seligmann and Jonathan Jacobs
# Copyright (C) 2012-2013 Bastian Kleineidam

from re import compile, MULTILINE, IGNORECASE, sub
from os.path import splitext
from ..scraper import _BasicScraper
from ..helpers import indirectStarter
from ..util import tagre


class SailorsunOrg(_BasicScraper):
    url = 'http://sailorsun.org/'
    stripUrl = url + '?p=%s'
    imageSearch = compile(tagre("img", "src", r'(http://sailorsun\.org/comics/[^"]+)'))
    prevSearch = compile(tagre("a", "href", r'(http://sailorsun\.org/\?p=\d+)', after="prev"))
    help = 'Index format: n (unpadded)'


class SamAndFuzzy(_BasicScraper):
    url = 'http://www.samandfuzzy.com/'
    stripUrl = 'http://samandfuzzy.com/%s'
    imageSearch = compile(r'(/comics/.+?)" alt')
    prevSearch = compile(r'"><a href="(.+?)"><img src="imgint/nav_prev.gif"')
    help = 'Index format: nnnn'


class SandraAndWoo(_BasicScraper):
    url = 'http://www.sandraandwoo.com/'
    stripUrl = url + '%s/'
    imageSearch = compile(tagre("img", "src", r'(http://www\.sandraandwoo\.com/comics/\d+-\d+-\d+-[^"]+)'))
    prevSearch = compile(tagre("a", "href", r'(http://www\.sandraandwoo\.com/\d+/\d+/\d+/[^"]+/)', after="prev"))
    help = 'Index format: yyyy/mm/dd/number-stripname'


class SarahZero(_BasicScraper):
    url = 'http://www.sarahzero.com/'
    stripUrl = url + 'sz_%s.html'
    imageSearch = compile(tagre("img", "src", r'(z_(?:spreads|decoy)/sz_[^"]+)'))
    prevSearch = compile(tagre("a", "href", r'(sz_\d+\.html)') + tagre("img", "src", r'z_site/sz_05_nav\.gif'))
    help = 'Index format: nnnn'


class ScaryGoRound(_BasicScraper):
    url = 'http://www.scarygoround.com/'
    stripUrl = url + '?date=%s'
    imageSearch = compile(tagre("img", "src", r'(strips/\d+\.png)'))
    prevSearch = compile(tagre("a", "href", r'(\?date=\d+)') + "Previous")
    help = 'Index format: n (unpadded)'


class ScenesFromAMultiverse(_BasicScraper):
    url = 'http://amultiverse.com/'
    stripUrl = url + '%s/'
    imageSearch = compile(tagre("img", "src", r'(http://amultiverse\.com/files/comics/[^"]+)'))
    prevSearch = compile(tagre("a", "href", r'(http://amultiverse\.com/\d+\d+/\d+/\d+/[^"]+)', after="prev"))
    help = 'Index format: yyyy/mm/dd/stripname'


class SchlockMercenary(_BasicScraper):
    url = 'http://www.schlockmercenary.com/'
    stripUrl = url + '%s'
    imageSearch = compile(tagre("img", "src", r'(http://static\.schlockmercenary\.com/comics/[^"]+)'))
    multipleImagesPerStrip = True
    prevSearch = compile(tagre("a", "href", r'(/\d+-\d+-\d+)', quote="'", after="nav-previous"))
    help = 'Index format: yyyy-mm-dd'


class SchoolBites(_BasicScraper):
    url = 'http://schoolbites.net/'
    stripUrl = url + 'd/%s.html'
    imageSearch = compile(tagre("img", "src", r'(http://cdn\.schoolbites\.net/comics/[^"]+)'))
    prevSearch = compile(tagre("a", "href", r'(http://schoolbites\.net/d/\d+\.html)', after="prev"))
    help = 'Index format: yyyymmdd'


class SequentialArt(_BasicScraper):
    url = 'http://www.collectedcurios.com/sequentialart.php'
    stripUrl = url + '?s=%s'
    imageSearch = compile(tagre("img", "src", r'([^"]+)', before="strip"))
    prevSearch = compile(tagre("a", "href", r'(/sequentialart\.php\?s=\d+)')
      + tagre("img", "src", "Nav_BackOne\.gif"))
    help = 'Index format: name'


class Sheldon(_BasicScraper):
    url = 'http://www.sheldoncomics.com/'
    stripUrl = url + 'archive/%s.html'
    imageSearch = compile(tagre("img", "src", r'(/strips/[^"]+)'))
    prevSearch = compile(tagre("a", "href", r'(/archive/\d+\.html)', after="sidenav-prev"))
    help = 'Index format: yymmdd'


class Shivae(_BasicScraper):
    url = 'http://shivae.net/'
    stripUrl = url + 'blog/%s/'
    imageSearch = compile(tagre("img", "src", r'(http://shivae\.net/files/comics/[^"]+)'))
    prevSearch = compile(tagre("a", "href", r'(http://shivae\.net/blog/[^"]+)', after="Previous"))
    help = 'Index format: yyyy/mm/dd/stripname'


# XXX disallowed by robots.txt
class _Shortpacked(_BasicScraper):
    url = 'http://www.shortpacked.com/'
    stripUrl = url + '%s/'
    imageSearch = compile(tagre("img", "src", r'(http://www\.shortpacked\.com/comics/[^"]+)'))
    prevSearch = compile(tagre("a", "href", r'(http://www\.shortpacked\.com/\d+/comic/[^"]+)', after="prev"))
    help = 'Index format: yyyy/comic/book-nn/mm-name1/name2'


class SinFest(_BasicScraper):
    name = 'KeenSpot/SinFest'
    url = 'http://www.sinfest.net/'
    stripUrl = url + 'archive_page.php?comicID=%s'
    imageSearch = compile(r'<img src=".+?(/comikaze/comics/.+?)"')
    prevSearch = compile(r'(/archive_page.php\?comicID=.+?)".+?prev_a')
    help = 'Index format: n (unpadded)'


class SkinDeep(_BasicScraper):
    url = 'http://www.skindeepcomic.com/'
    stripUrl = url + 'archive/%s/'
    imageSearch = compile(r'<span class="webcomic-object[^>]*><img src="([^"]*)"')
    prevSearch = compile(tagre("a", "href", r'([^"]+)', after="previous-webcomic-link"))
    help = 'Index format: custom'


class SlightlyDamned(_BasicScraper):
    url = 'http://www.sdamned.com/'
    stripUrl = url + '%s/'
    imageSearch = compile(tagre("img", "src", r'(http://www\.sdamned\.com/comics/[^"]+)'))
    prevSearch = compile(tagre("a", "href", r'(http://www\.sdamned\.com/[^"]+)', after="prev"))
    help = 'Index format: yyyy/mm/number'


class SluggyFreelance(_BasicScraper):
    url = 'http://www.sluggy.com/'
    stripUrl = url + 'comics/archives/daily/%s'
    imageSearch = compile(r'<img src="(/images/comics/.+?)"')
    prevSearch = compile(r'<a href="(.+?)"[^>]+?><span class="ui-icon ui-icon-seek-prev">')
    help = 'Index format: yymmdd'


class SodiumEyes(_BasicScraper):
    url = 'http://sodiumeyes.com/'
    stripUrl = url + '%s/'
    imageSearch = compile(tagre("img", "src", r'(http://sodiumeyes\.com/comic/[^ ]+)', quote=""))
    prevSearch = compile(tagre("a", "href", r'(http://sodiumeyes\.com/[^"]+)', after="prev"))
    help = 'Index format: yyyy/mm/dd/stripname'


class Sorcery101(_BasicScraper):
    url = 'http://www.sorcery101.net/sorcery-101/'
    stripUrl = url + '%s/'
    imageSearch = compile(tagre("img", "src", r'(http://www\.sorcery101\.net/wp-content/uploads/\d+/\d+/[^"]+)'))
    prevSearch = compile(tagre("a", "href", r'(http://www\.sorcery101\.net/sorcery-101/[^"]+)', after="previous-"))
    help = 'Index format: stripname'


class SpaceTrawler(_BasicScraper):
    url = 'http://spacetrawler.com/'
    stripUrl = url + '%s/'
    imageSearch = compile(tagre("img", "src", r'(http://spacetrawler\.com/comics/[^"]+)'))
    prevSearch = compile(tagre("a", "href", r'(http://spacetrawler\.com/\d+/\d+/\d+/[^"]+)', after="navi-prev"))
    help = 'Index format: yyyy/mm/dd/stripname'


class SpareParts(_BasicScraper):
    baseUrl = 'http://www.sparepartscomics.com/'
    url = baseUrl + 'comics/?date=20080328'
    stripUrl = baseUrl + 'comics/index.php?date=%s'
    imageSearch = compile(tagre("img", "src", r'(http://www\.sparepartscomics\.com/comics/[^"]+)'))
    prevSearch = compile(tagre("a", "href", r'(index\.php\?date=\d+)', quote="'") + "Previous Comic")
    help = 'Index format: yyyymmdd'


class Spinnerette(_BasicScraper):
    url = 'http://www.spinnyverse.com/'
    stripUrl = url + '%s/'
    imageSearch = compile(tagre("img", "src", r'(http://www\.spinnyverse\.com/comics/[^"]+)'))
    prevSearch = compile(tagre("a", "href", r'(http://www\.spinnyverse\.com/[^"]+)', before="Previous Comic"))
    help = 'Index format: number'


class SPQRBlues(_BasicScraper):
    url = 'http://spqrblues.com/IV/'
    stripUrl = url + '?p=%s'
    imageSearch = compile(tagre("img", "src", r'(http://spqrblues\.com/IV/comics/\d+\.png)'))
    prevSearch = compile(tagre("a", "href", r'(http://spqrblues\.com/IV/\?p=\d+)', after="prev"))
    help = 'Index format: number'


# XXX disallowed by robots.txt
class _StationV3(_BasicScraper):
    url = 'http://www.stationv3.com/'
    stripUrl = url + 'd/%s.html'
    imageSearch = compile(tagre("img", "src", r'(http://www\.stationv3\.com/comics/[^"]+)'))
    prevSearch = compile(tagre("a", "href", r'(http://www\.stationv3\.com/d/\d+\.html)') +
      tagre("img", "src", r'http://www\.stationv3\.com/images/previous\.gif'))
    help = 'Index format: yyyymmdd'


class Stubble(_BasicScraper):
    url = 'http://stubblecomics.com/'
    stripUrl = url + '?p=%s'
    imageSearch = compile(tagre("img", "src", r'(http://stubblecomics\.com/comics/[^"]+)'))
    prevSearch = compile(tagre("a", "href", r'(http://stubblecomics\.com/\?p=\d+)', after="navi-prev"))
    help = 'Index format: number'


class StrawberryDeathCake(_BasicScraper):
    url = 'http://strawberrydeathcake.com/'
    stripUrl = url + 'archive/%s/'
    imageSearch = compile(tagre("img", "src", r'(http://strawberrydeathcake\.com/wp-content/webcomic/[^"]+)'))
    prevSearch = compile(tagre("a", "href", r'(http://strawberrydeathcake\.com/archive/[^"]+)', after="previous"))
    help = 'Index format: stripname'


class SuburbanTribe(_BasicScraper):
    url = 'http://www.pixelwhip.com/'
    stripUrl = url + '?p=%s'
    imageSearch = compile(tagre("img", "src", r'(http://www\.pixelwhip\.com/comics/[^"]+)'))
    prevSearch = compile(tagre("a", "href", r'(http://www\.pixelwhip\.com/\?p=\d+)', after="prev"))
    help = 'Index format: nnnn'


class SomethingPositive(_BasicScraper):
    url = 'http://www.somethingpositive.net/'
    stripUrl = url + 'sp%s.shtml'
    imageSearch = compile(tagre("img", "src", r'(sp\d+\.png)'))
    prevSearch = compile(tagre("a", "href", r'(sp\d+\.shtml)') +
      "(?:" + tagre("img", "src", r'images/previous\.gif') + "|Previous)")
    help = 'Index format: mmddyyyy'


class SexyLosers(_BasicScraper):
    adult = True
    url = 'http://www.sexylosers.com/'
    stripUrl = url + '%s.html'
    imageSearch = compile(r'<img src\s*=\s*"\s*(comics/[\w\.]+?)"', IGNORECASE)
    prevSearch = compile(r'<a href="(/\d{3}\.\w+?)"><font color = FFAAAA><<', IGNORECASE)
    help = 'Index format: nnn'
    starter = indirectStarter(url,
                              compile(r'SEXY LOSERS <A HREF="(.+?)">Latest SL Comic \(#\d+\)</A>', IGNORECASE))

    @classmethod
    def namer(cls, imageUrl, pageUrl):
        index = pageUrl.split('/')[-1].split('.')[0]
        title = imageUrl.split('/')[-1].split('.')[0]
        return index + '-' + title


class StarCrossdDestiny(_BasicScraper):
    url = 'http://www.starcrossd.net/comic.html'
    stripUrl = 'http://www.starcrossd.net/archives/%s.html'
    imageSearch = compile(tagre("img", "src", r'(http://www\.starcrossd\.net/(?:ch1|strips|book2)/[^"]+)'))
    prevSearch = compile(r'<a href="(http://www\.starcrossd\.net/(?:ch1/)?archives/\d+\.html)"[^>]*"[^"]*"[^>]*>prev', IGNORECASE)
    help = 'Index format: nnnnnnnn'

    @classmethod
    def namer(cls, imageUrl, pageUrl):
        if imageUrl.find('ch1') == -1:
            # At first all images were stored in a strips/ directory but that was changed with the introduction of book2
            imageUrl = sub('(?:strips)|(?:images)','book1',imageUrl)
        elif not imageUrl.find('strips') == -1:
            imageUrl = imageUrl.replace('strips/','')
        directory, filename = imageUrl.split('/')[-2:]
        filename, extension = splitext(filename)
        return directory + '-' + filename


class Spamusement(_BasicScraper):
    url = 'http://spamusement.com/'
    stripUrl = url + 'index.php/comics/view/%s'
    imageSearch = compile(r'<img src="(http://spamusement.com/gfx/\d+\..+?)"', IGNORECASE)
    prevSearch = compile(r'<a href="(http://spamusement.com/index.php/comics/view/.+?)">', IGNORECASE)
    help = 'Index format: n (unpadded)'
    starter = indirectStarter(url, prevSearch)


# XXX disallowed by robots.txt
class _StrangeCandy(_BasicScraper):
    url = 'http://www.strangecandy.net/'
    stripUrl = url + 'd/%s.html'
    imageSearch = compile(tagre("img", "src", r'(/comics/\d+\.jpg)'))
    prevSearch = compile(tagre("a", "href", r'(/d/\d+\.html)') + tagre("img", "alt", "Previous comic"))
    help = 'Index format: yyyyddmm'


class SMBC(_BasicScraper):
    url = 'http://www.smbc-comics.com/'
    stripUrl = url + 'index.php?db=comics&id=%s'
    imageSearch = compile(r'<img src=\'(.+?\d{8}.\w{1,4})\'>')
    prevSearch = compile(r'131,13,216,84"\n\s+href="(.+?)#comic"\n>', MULTILINE)
    help = 'Index format: nnnn'


class SupernormalStep(_BasicScraper):
    url = 'http://supernormalstep.com/'
    stripUrl = url + '?p=%s'
    imageSearch = compile(tagre("img", "src", r'(http://supernormalstep\.com/comics/[^"]+)'))
    prevSearch = compile(tagre("a", "href", r'(http://supernormalstep\.com/\?p=\d+)', after="prev"))
    help = 'Index format: number'
Updated copyright for all source files. 2012-06-20 20:41:04 +00:00			`# -- coding: iso-8859-1 --`
			`# Copyright (C) 2004-2005 Tristan Seligmann and Jonathan Jacobs`
Rename latestUrl in url 2013-02-05 18:51:46 +00:00			`# Copyright (C) 2012-2013 Bastian Kleineidam`
Fix some comics. 2012-11-21 20:57:26 +00:00
Initial commit to Github. 2012-06-20 19:58:13 +00:00			`from re import compile, MULTILINE, IGNORECASE, sub`
			`from os.path import splitext`
A lot of refactoring. 2012-10-11 10:03:12 +00:00			`from ..scraper import _BasicScraper`
Fix some comics. 2012-11-21 20:57:26 +00:00			`from ..helpers import indirectStarter`
Updated documentation and fix some comics. 2012-11-20 17:53:53 +00:00			`from ..util import tagre`
Initial commit to Github. 2012-06-20 19:58:13 +00:00

			`class SailorsunOrg(_BasicScraper):`
Always have an url attribute in comic scrapers. 2013-02-04 20:00:26 +00:00			`url = 'http://sailorsun.org/'`
			`stripUrl = url + '?p=%s'`
Fix some comics. 2012-11-21 20:57:26 +00:00			`imageSearch = compile(tagre("img", "src", r'(http://sailorsun\.org/comics/[^"]+)'))`
			`prevSearch = compile(tagre("a", "href", r'(http://sailorsun\.org/\?p=\d+)', after="prev"))`
Initial commit to Github. 2012-06-20 19:58:13 +00:00			`help = 'Index format: n (unpadded)'`


			`class SamAndFuzzy(_BasicScraper):`
Always have an url attribute in comic scrapers. 2013-02-04 20:00:26 +00:00			`url = 'http://www.samandfuzzy.com/'`
Rename imageUrl to stripUrl. 2012-11-13 18:10:19 +00:00			`stripUrl = 'http://samandfuzzy.com/%s'`
Initial commit to Github. 2012-06-20 19:58:13 +00:00			`imageSearch = compile(r'(/comics/.+?)" alt')`
			`prevSearch = compile(r'"><a href="(.+?)"><img src="imgint/nav_prev.gif"')`
			`help = 'Index format: nnnn'`


Add SandraAndWoo, SupernormalStep 2013-02-13 16:53:11 +00:00			`class SandraAndWoo(_BasicScraper):`
			`url = 'http://www.sandraandwoo.com/'`
			`stripUrl = url + '%s/'`
			`imageSearch = compile(tagre("img", "src", r'(http://www\.sandraandwoo\.com/comics/\d+-\d+-\d+-[^"]+)'))`
			`prevSearch = compile(tagre("a", "href", r'(http://www\.sandraandwoo\.com/\d+/\d+/\d+/[^"]+/)', after="prev"))`
			`help = 'Index format: yyyy/mm/dd/number-stripname'`


Initial commit to Github. 2012-06-20 19:58:13 +00:00			`class SarahZero(_BasicScraper):`
Always have an url attribute in comic scrapers. 2013-02-04 20:00:26 +00:00			`url = 'http://www.sarahzero.com/'`
			`stripUrl = url + 'sz_%s.html'`
Fix comics. 2012-12-04 06:02:40 +00:00			`imageSearch = compile(tagre("img", "src", r'(z_(?:spreads\|decoy)/sz_[^"]+)'))`
Fix some comics. 2012-11-21 20:57:26 +00:00			`prevSearch = compile(tagre("a", "href", r'(sz_\d+\.html)') + tagre("img", "src", r'z_site/sz_05_nav\.gif'))`
Initial commit to Github. 2012-06-20 19:58:13 +00:00			`help = 'Index format: nnnn'`


			`class ScaryGoRound(_BasicScraper):`
Always have an url attribute in comic scrapers. 2013-02-04 20:00:26 +00:00			`url = 'http://www.scarygoround.com/'`
			`stripUrl = url + '?date=%s'`
Fix some comics. 2012-11-21 20:57:26 +00:00			`imageSearch = compile(tagre("img", "src", r'(strips/\d+\.png)'))`
			`prevSearch = compile(tagre("a", "href", r'(\?date=\d+)') + "Previous")`
Initial commit to Github. 2012-06-20 19:58:13 +00:00			`help = 'Index format: n (unpadded)'`


Added some comics. 2013-02-06 21:08:36 +00:00			`class ScenesFromAMultiverse(_BasicScraper):`
			`url = 'http://amultiverse.com/'`
			`stripUrl = url + '%s/'`
			`imageSearch = compile(tagre("img", "src", r'(http://amultiverse\.com/files/comics/[^"]+)'))`
			`prevSearch = compile(tagre("a", "href", r'(http://amultiverse\.com/\d+\d+/\d+/\d+/[^"]+)', after="prev"))`
			`help = 'Index format: yyyy/mm/dd/stripname'`


Updated documentation and fix some comics. 2012-11-20 17:53:53 +00:00			`class SchlockMercenary(_BasicScraper):`
Always have an url attribute in comic scrapers. 2013-02-04 20:00:26 +00:00			`url = 'http://www.schlockmercenary.com/'`
			`stripUrl = url + '%s'`
Fix some comics. 2012-11-21 20:57:26 +00:00			`imageSearch = compile(tagre("img", "src", r'(http://static\.schlockmercenary\.com/comics/[^"]+)'))`
Fix comics. 2012-12-04 06:02:40 +00:00			`multipleImagesPerStrip = True`
			`prevSearch = compile(tagre("a", "href", r'(/\d+-\d+-\d+)', quote="'", after="nav-previous"))`
Fix some comics. 2012-11-21 20:57:26 +00:00			`help = 'Index format: yyyy-mm-dd'`
Updated documentation and fix some comics. 2012-11-20 17:53:53 +00:00
Initial commit to Github. 2012-06-20 19:58:13 +00:00
			`class SchoolBites(_BasicScraper):`
Always have an url attribute in comic scrapers. 2013-02-04 20:00:26 +00:00			`url = 'http://schoolbites.net/'`
			`stripUrl = url + 'd/%s.html'`
Fix some comics. 2012-11-21 20:57:26 +00:00			`imageSearch = compile(tagre("img", "src", r'(http://cdn\.schoolbites\.net/comics/[^"]+)'))`
			`prevSearch = compile(tagre("a", "href", r'(http://schoolbites\.net/d/\d+\.html)', after="prev"))`
Initial commit to Github. 2012-06-20 19:58:13 +00:00			`help = 'Index format: yyyymmdd'`


Add SequentialArt comic. 2013-01-29 20:23:32 +00:00			`class SequentialArt(_BasicScraper):`
Always have an url attribute in comic scrapers. 2013-02-04 20:00:26 +00:00			`url = 'http://www.collectedcurios.com/sequentialart.php'`
			`stripUrl = url + '?s=%s'`
Add SequentialArt comic. 2013-01-29 20:23:32 +00:00			`imageSearch = compile(tagre("img", "src", r'([^"]+)', before="strip"))`
			`prevSearch = compile(tagre("a", "href", r'(/sequentialart\.php\?s=\d+)')`
			`+ tagre("img", "src", "Nav_BackOne\.gif"))`
			`help = 'Index format: name'`


Updated documentation and fix some comics. 2012-11-20 17:53:53 +00:00			`class Sheldon(_BasicScraper):`
Always have an url attribute in comic scrapers. 2013-02-04 20:00:26 +00:00			`url = 'http://www.sheldoncomics.com/'`
			`stripUrl = url + 'archive/%s.html'`
Fix some comics. 2012-11-21 20:57:26 +00:00			`imageSearch = compile(tagre("img", "src", r'(/strips/[^"]+)'))`
Fix comics, improve tests, use python-requests. 2012-11-26 17:44:31 +00:00			`prevSearch = compile(tagre("a", "href", r'(/archive/\d+\.html)', after="sidenav-prev"))`
Fix some comics. 2012-11-21 20:57:26 +00:00			`help = 'Index format: yymmdd'`
Updated documentation and fix some comics. 2012-11-20 17:53:53 +00:00

Added comics. 2012-12-08 20:30:51 +00:00			`class Shivae(_BasicScraper):`
Always have an url attribute in comic scrapers. 2013-02-04 20:00:26 +00:00			`url = 'http://shivae.net/'`
			`stripUrl = url + 'blog/%s/'`
Added comics. 2012-12-08 20:30:51 +00:00			`imageSearch = compile(tagre("img", "src", r'(http://shivae\.net/files/comics/[^"]+)'))`
			`prevSearch = compile(tagre("a", "href", r'(http://shivae\.net/blog/[^"]+)', after="Previous"))`
			`help = 'Index format: yyyy/mm/dd/stripname'`


Various fixes and additions. 2012-12-12 16:41:29 +00:00			`# XXX disallowed by robots.txt`
			`class _Shortpacked(_BasicScraper):`
Always have an url attribute in comic scrapers. 2013-02-04 20:00:26 +00:00			`url = 'http://www.shortpacked.com/'`
			`stripUrl = url + '%s/'`
Fix some comics. 2012-11-21 20:57:26 +00:00			`imageSearch = compile(tagre("img", "src", r'(http://www\.shortpacked\.com/comics/[^"]+)'))`
			`prevSearch = compile(tagre("a", "href", r'(http://www\.shortpacked\.com/\d+/comic/[^"]+)', after="prev"))`
Fix some comics. 2012-12-02 17:35:06 +00:00			`help = 'Index format: yyyy/comic/book-nn/mm-name1/name2'`
Updated documentation and fix some comics. 2012-11-20 17:53:53 +00:00

Initial commit to Github. 2012-06-20 19:58:13 +00:00			`class SinFest(_BasicScraper):`
			`name = 'KeenSpot/SinFest'`
Always have an url attribute in comic scrapers. 2013-02-04 20:00:26 +00:00			`url = 'http://www.sinfest.net/'`
			`stripUrl = url + 'archive_page.php?comicID=%s'`
Initial commit to Github. 2012-06-20 19:58:13 +00:00			`imageSearch = compile(r'<img src=".+?(/comikaze/comics/.+?)"')`
			`prevSearch = compile(r'(/archive_page.php\?comicID=.+?)".+?prev_a')`
			`help = 'Index format: n (unpadded)'`


Add SkinDeep. Filenames for this are all over the place :( 2013-02-07 22:02:54 +00:00			`class SkinDeep(_BasicScraper):`
			`url = 'http://www.skindeepcomic.com/'`
			`stripUrl = url + 'archive/%s/'`
			`imageSearch = compile(r'<span class="webcomic-object[^>]><img src="([^"])"')`
			`prevSearch = compile(tagre("a", "href", r'([^"]+)', after="previous-webcomic-link"))`
			`help = 'Index format: custom'`


Initial commit to Github. 2012-06-20 19:58:13 +00:00			`class SlightlyDamned(_BasicScraper):`
Always have an url attribute in comic scrapers. 2013-02-04 20:00:26 +00:00			`url = 'http://www.sdamned.com/'`
			`stripUrl = url + '%s/'`
Fix some comics. 2012-11-21 20:57:26 +00:00			`imageSearch = compile(tagre("img", "src", r'(http://www\.sdamned\.com/comics/[^"]+)'))`
			`prevSearch = compile(tagre("a", "href", r'(http://www\.sdamned\.com/[^"]+)', after="prev"))`
			`help = 'Index format: yyyy/mm/number'`
Initial commit to Github. 2012-06-20 19:58:13 +00:00

			`class SluggyFreelance(_BasicScraper):`
Always have an url attribute in comic scrapers. 2013-02-04 20:00:26 +00:00			`url = 'http://www.sluggy.com/'`
			`stripUrl = url + 'comics/archives/daily/%s'`
Initial commit to Github. 2012-06-20 19:58:13 +00:00			`imageSearch = compile(r'<img src="(/images/comics/.+?)"')`
			`prevSearch = compile(r'<a href="(.+?)"[^>]+?><span class="ui-icon ui-icon-seek-prev">')`
			`help = 'Index format: yymmdd'`


			`class SodiumEyes(_BasicScraper):`
Always have an url attribute in comic scrapers. 2013-02-04 20:00:26 +00:00			`url = 'http://sodiumeyes.com/'`
			`stripUrl = url + '%s/'`
Fix comics. 2012-12-04 06:02:40 +00:00			`imageSearch = compile(tagre("img", "src", r'(http://sodiumeyes\.com/comic/[^ ]+)', quote=""))`
Fix some comics. 2012-11-21 20:57:26 +00:00			`prevSearch = compile(tagre("a", "href", r'(http://sodiumeyes\.com/[^"]+)', after="prev"))`
			`help = 'Index format: yyyy/mm/dd/stripname'`
Initial commit to Github. 2012-06-20 19:58:13 +00:00

Added comics. 2012-12-08 20:30:51 +00:00			`class Sorcery101(_BasicScraper):`
Fix some comics. 2013-02-27 18:40:54 +00:00			`url = 'http://www.sorcery101.net/sorcery-101/'`
			`stripUrl = url + '%s/'`
			`imageSearch = compile(tagre("img", "src", r'(http://www\.sorcery101\.net/wp-content/uploads/\d+/\d+/[^"]+)'))`
			`prevSearch = compile(tagre("a", "href", r'(http://www\.sorcery101\.net/sorcery-101/[^"]+)', after="previous-"))`
Added comics. 2012-12-08 20:30:51 +00:00			`help = 'Index format: stripname'`


Added some comics. 2013-02-06 21:08:36 +00:00			`class SpaceTrawler(_BasicScraper):`
			`url = 'http://spacetrawler.com/'`
			`stripUrl = url + '%s/'`
			`imageSearch = compile(tagre("img", "src", r'(http://spacetrawler\.com/comics/[^"]+)'))`
			`prevSearch = compile(tagre("a", "href", r'(http://spacetrawler\.com/\d+/\d+/\d+/[^"]+)', after="navi-prev"))`
			`help = 'Index format: yyyy/mm/dd/stripname'`


Initial commit to Github. 2012-06-20 19:58:13 +00:00			`class SpareParts(_BasicScraper):`
Fix some comics. 2012-11-21 20:57:26 +00:00			`baseUrl = 'http://www.sparepartscomics.com/'`
Always have an url attribute in comic scrapers. 2013-02-04 20:00:26 +00:00			`url = baseUrl + 'comics/?date=20080328'`
Fix comics. 2012-12-04 06:02:40 +00:00			`stripUrl = baseUrl + 'comics/index.php?date=%s'`
			`imageSearch = compile(tagre("img", "src", r'(http://www\.sparepartscomics\.com/comics/[^"]+)'))`
			`prevSearch = compile(tagre("a", "href", r'(index\.php\?date=\d+)', quote="'") + "Previous Comic")`
Updated documentation and fix some comics. 2012-11-20 17:53:53 +00:00			`help = 'Index format: yyyymmdd'`

Initial commit to Github. 2012-06-20 19:58:13 +00:00
Add Spinnerette comic. 2013-01-29 20:52:26 +00:00			`class Spinnerette(_BasicScraper):`
Always have an url attribute in comic scrapers. 2013-02-04 20:00:26 +00:00			`url = 'http://www.spinnyverse.com/'`
			`stripUrl = url + '%s/'`
Add Spinnerette comic. 2013-01-29 20:52:26 +00:00			`imageSearch = compile(tagre("img", "src", r'(http://www\.spinnyverse\.com/comics/[^"]+)'))`
Fix Spinnerette. The old expression was matching "Previous issue" first and skipping all comics. 2013-02-07 21:29:51 +00:00			`prevSearch = compile(tagre("a", "href", r'(http://www\.spinnyverse\.com/[^"]+)', before="Previous Comic"))`
Add Spinnerette comic. 2013-01-29 20:52:26 +00:00			`help = 'Index format: number'`


Added comics. 2012-12-08 20:30:51 +00:00			`class SPQRBlues(_BasicScraper):`
Always have an url attribute in comic scrapers. 2013-02-04 20:00:26 +00:00			`url = 'http://spqrblues.com/IV/'`
			`stripUrl = url + '?p=%s'`
Added comics. 2012-12-08 20:30:51 +00:00			`imageSearch = compile(tagre("img", "src", r'(http://spqrblues\.com/IV/comics/\d+\.png)'))`
			`prevSearch = compile(tagre("a", "href", r'(http://spqrblues\.com/IV/\?p=\d+)', after="prev"))`
			`help = 'Index format: number'`


Various comics are fixed. 2012-12-13 20:05:27 +00:00			`# XXX disallowed by robots.txt`
			`class _StationV3(_BasicScraper):`
Always have an url attribute in comic scrapers. 2013-02-04 20:00:26 +00:00			`url = 'http://www.stationv3.com/'`
			`stripUrl = url + 'd/%s.html'`
Added comics. 2012-12-08 20:30:51 +00:00			`imageSearch = compile(tagre("img", "src", r'(http://www\.stationv3\.com/comics/[^"]+)'))`
			`prevSearch = compile(tagre("a", "href", r'(http://www\.stationv3\.com/d/\d+\.html)') +`
			`tagre("img", "src", r'http://www\.stationv3\.com/images/previous\.gif'))`
			`help = 'Index format: yyyymmdd'`


Initial commit to Github. 2012-06-20 19:58:13 +00:00			`class Stubble(_BasicScraper):`
Always have an url attribute in comic scrapers. 2013-02-04 20:00:26 +00:00			`url = 'http://stubblecomics.com/'`
			`stripUrl = url + '?p=%s'`
Fix some comics. 2012-11-21 20:57:26 +00:00			`imageSearch = compile(tagre("img", "src", r'(http://stubblecomics\.com/comics/[^"]+)'))`
			`prevSearch = compile(tagre("a", "href", r'(http://stubblecomics\.com/\?p=\d+)', after="navi-prev"))`
			`help = 'Index format: number'`
Initial commit to Github. 2012-06-20 19:58:13 +00:00

			`class StrawberryDeathCake(_BasicScraper):`
Always have an url attribute in comic scrapers. 2013-02-04 20:00:26 +00:00			`url = 'http://strawberrydeathcake.com/'`
			`stripUrl = url + 'archive/%s/'`
Fix comics. 2012-12-04 06:02:40 +00:00			`imageSearch = compile(tagre("img", "src", r'(http://strawberrydeathcake\.com/wp-content/webcomic/[^"]+)'))`
Fix some comics. 2012-11-21 20:57:26 +00:00			`prevSearch = compile(tagre("a", "href", r'(http://strawberrydeathcake\.com/archive/[^"]+)', after="previous"))`
			`help = 'Index format: stripname'`
Initial commit to Github. 2012-06-20 19:58:13 +00:00

			`class SuburbanTribe(_BasicScraper):`
Always have an url attribute in comic scrapers. 2013-02-04 20:00:26 +00:00			`url = 'http://www.pixelwhip.com/'`
			`stripUrl = url + '?p=%s'`
Fix some comics. 2012-11-21 20:57:26 +00:00			`imageSearch = compile(tagre("img", "src", r'(http://www\.pixelwhip\.com/comics/[^"]+)'))`
			`prevSearch = compile(tagre("a", "href", r'(http://www\.pixelwhip\.com/\?p=\d+)', after="prev"))`
Initial commit to Github. 2012-06-20 19:58:13 +00:00			`help = 'Index format: nnnn'`


			`class SomethingPositive(_BasicScraper):`
Always have an url attribute in comic scrapers. 2013-02-04 20:00:26 +00:00			`url = 'http://www.somethingpositive.net/'`
			`stripUrl = url + 'sp%s.shtml'`
Fix some comics. 2012-11-21 20:57:26 +00:00			`imageSearch = compile(tagre("img", "src", r'(sp\d+\.png)'))`
Various fixes and additions. 2012-12-12 16:41:29 +00:00			`prevSearch = compile(tagre("a", "href", r'(sp\d+\.shtml)') +`
Fix comics. 2012-12-04 06:02:40 +00:00			`"(?:" + tagre("img", "src", r'images/previous\.gif') + "\|Previous)")`
Initial commit to Github. 2012-06-20 19:58:13 +00:00			`help = 'Index format: mmddyyyy'`


			`class SexyLosers(_BasicScraper):`
Cleanup 2012-12-09 19:15:22 +00:00			`adult = True`
Always have an url attribute in comic scrapers. 2013-02-04 20:00:26 +00:00			`url = 'http://www.sexylosers.com/'`
			`stripUrl = url + '%s.html'`
Initial commit to Github. 2012-06-20 19:58:13 +00:00			`imageSearch = compile(r'<img src\s=\s"\s*(comics/[\w\.]+?)"', IGNORECASE)`
			`prevSearch = compile(r'<a href="(/\d{3}\.\w+?)"><font color = FFAAAA><<', IGNORECASE)`
			`help = 'Index format: nnn'`
Always have an url attribute in comic scrapers. 2013-02-04 20:00:26 +00:00			`starter = indirectStarter(url,`
Initial commit to Github. 2012-06-20 19:58:13 +00:00			`compile(r'SEXY LOSERS <A HREF="(.+?)">Latest SL Comic \(#\d+\)</A>', IGNORECASE))`

			`@classmethod`
			`def namer(cls, imageUrl, pageUrl):`
			`index = pageUrl.split('/')[-1].split('.')[0]`
			`title = imageUrl.split('/')[-1].split('.')[0]`
			`return index + '-' + title`


			`class StarCrossdDestiny(_BasicScraper):`
Always have an url attribute in comic scrapers. 2013-02-04 20:00:26 +00:00			`url = 'http://www.starcrossd.net/comic.html'`
Rename imageUrl to stripUrl. 2012-11-13 18:10:19 +00:00			`stripUrl = 'http://www.starcrossd.net/archives/%s.html'`
Fix comics. 2012-12-04 06:02:40 +00:00			`imageSearch = compile(tagre("img", "src", r'(http://www\.starcrossd\.net/(?:ch1\|strips\|book2)/[^"]+)'))`
Initial commit to Github. 2012-06-20 19:58:13 +00:00			`prevSearch = compile(r'<a href="(http://www\.starcrossd\.net/(?:ch1/)?archives/\d+\.html)"[^>]"[^"]"[^>]*>prev', IGNORECASE)`
			`help = 'Index format: nnnnnnnn'`

			`@classmethod`
			`def namer(cls, imageUrl, pageUrl):`
			`if imageUrl.find('ch1') == -1:`
			`# At first all images were stored in a strips/ directory but that was changed with the introduction of book2`
			`imageUrl = sub('(?:strips)\|(?:images)','book1',imageUrl)`
			`elif not imageUrl.find('strips') == -1:`
			`imageUrl = imageUrl.replace('strips/','')`
			`directory, filename = imageUrl.split('/')[-2:]`
			`filename, extension = splitext(filename)`
			`return directory + '-' + filename`


			`class Spamusement(_BasicScraper):`
Always have an url attribute in comic scrapers. 2013-02-04 20:00:26 +00:00			`url = 'http://spamusement.com/'`
			`stripUrl = url + 'index.php/comics/view/%s'`
Initial commit to Github. 2012-06-20 19:58:13 +00:00			`imageSearch = compile(r'<img src="(http://spamusement.com/gfx/\d+\..+?)"', IGNORECASE)`
			`prevSearch = compile(r'<a href="(http://spamusement.com/index.php/comics/view/.+?)">', IGNORECASE)`
			`help = 'Index format: n (unpadded)'`
Always have an url attribute in comic scrapers. 2013-02-04 20:00:26 +00:00			`starter = indirectStarter(url, prevSearch)`
Initial commit to Github. 2012-06-20 19:58:13 +00:00

Various comics are fixed. 2012-12-13 20:05:27 +00:00			`# XXX disallowed by robots.txt`
			`class _StrangeCandy(_BasicScraper):`
Always have an url attribute in comic scrapers. 2013-02-04 20:00:26 +00:00			`url = 'http://www.strangecandy.net/'`
			`stripUrl = url + 'd/%s.html'`
Fix some comics. 2012-11-21 20:57:26 +00:00			`imageSearch = compile(tagre("img", "src", r'(/comics/\d+\.jpg)'))`
			`prevSearch = compile(tagre("a", "href", r'(/d/\d+\.html)') + tagre("img", "alt", "Previous comic"))`
Initial commit to Github. 2012-06-20 19:58:13 +00:00			`help = 'Index format: yyyyddmm'`


			`class SMBC(_BasicScraper):`
Always have an url attribute in comic scrapers. 2013-02-04 20:00:26 +00:00			`url = 'http://www.smbc-comics.com/'`
			`stripUrl = url + 'index.php?db=comics&id=%s'`
Initial commit to Github. 2012-06-20 19:58:13 +00:00			`imageSearch = compile(r'<img src=\'(.+?\d{8}.\w{1,4})\'>')`
			`prevSearch = compile(r'131,13,216,84"\n\s+href="(.+?)#comic"\n>', MULTILINE)`
			`help = 'Index format: nnnn'`
Add SandraAndWoo, SupernormalStep 2013-02-13 16:53:11 +00:00

			`class SupernormalStep(_BasicScraper):`
			`url = 'http://supernormalstep.com/'`
			`stripUrl = url + '?p=%s'`
			`imageSearch = compile(tagre("img", "src", r'(http://supernormalstep\.com/comics/[^"]+)'))`
			`prevSearch = compile(tagre("a", "href", r'(http://supernormalstep\.com/\?p=\d+)', after="prev"))`
			`help = 'Index format: number'`