X-Git-Url: http://www.chiark.greenend.org.uk/ucgi/~yarrgweb/git?a=blobdiff_plain;f=yoweb-scrape;h=2d968921ab1eed4023fea7344ce90951f2c7c21d;hb=31cc3a65d81a3b9edd09372950ffbf3da5ab283b;hp=462bd01c0366a42e877d8b47b8c4dade0a9c34ec;hpb=2b1646498eea5609775d17e121937feea4aa1196;p=ypp-sc-tools.db-test.git diff --git a/yoweb-scrape b/yoweb-scrape index 462bd01..2d96892 100755 --- a/yoweb-scrape +++ b/yoweb-scrape @@ -33,7 +33,7 @@ max_pirate_namelen = 12 def debug(m): if opts.debug: - print >>sys.stderr, m + print m class Fetcher: def __init__(self, ocean, cachedir): @@ -70,7 +70,7 @@ class Fetcher: ages.append(age) return ages - def _rate_limit_cache_clean(self, now): + def need_wait(self, now): ages = self._cache_scan(now) ages.sort() debug('Fetcher ages ' + `ages`) @@ -83,6 +83,10 @@ class Fetcher: need_wait = max(need_wait, min_age - age) min_age += 3 min_age *= 1.25 + return need_wait + + def _rate_limit_cache_clean(self, now): + need_wait = self.need_wait(now) if need_wait > 0: debug('Fetcher wait %d' % need_wait) time.sleep(need_wait) @@ -280,7 +284,7 @@ class CrewInfo(SomethingSoupInfo): crew_rank_re = regexp.compile('/yoweb/images/crew') for row in tbl.contents: # findAll(recurse=False) - if isinstance(row, unicode): + if isinstance(row,basestring): continue is_rank = row.find('img', attrs={'src': crew_rank_re}) @@ -327,6 +331,7 @@ class StandingsTable: if not isinstance(puzzle,list): puzzle = [puzzle] try: standing = max([pi.standings[p] for p in puzzle]) except KeyError: return '?' + if isinstance(standing,basestring): return standing if not standing: return '' s = '' if self._cw > 4: @@ -391,14 +396,27 @@ class PirateAboard: # pa.last_time # pa.last_event # pa.gunner - def __init__(pa, v, time, event): + # pa.last_chat_time + # pa.last_chat_chan + # pa.pi + + def __init__(pa, pn, v, time, event): + pa.pn = pn pa.v = v pa.last_time = time pa.last_event = event + pa.last_chat_time = None + pa.last_chat_chan = None pa.gunner = False + pa.pi = None -class ShipCrewTracker: - def __init__(self, myself_pi): + def pirate_info(self): + if not pa.pi and not fetcher.need_wait(time.time): + pa.pi = PirateInfo(pa.name, 3600) + return pa.pi + +class ChatLogTracker: + def __init__(self, myself_pi, logfn): self._pl = {} # self._pl['Pirate'] = self._vl = {} # self._vl['Vessel']['Pirate'] = PirateAboard # self._vl['Vessel']['#lastaboard'] @@ -406,10 +424,13 @@ class ShipCrewTracker: self._vessel = None # self._vl[self._vessel] self._date = None self._myself = myself_pi - self._need_redisplay = False + self.need_redisplay = False + self._f = file(logfn) + self._lbuf = '' + self._progress = [0, os.fstat(self._f.fileno()).st_size] def _refresh(self): - self._need_redisplay = True + self.need_redisplay = True def _onboard_event(self,timestamp,pirate,event): try: pa = self._pl[pirate] @@ -419,22 +440,36 @@ class ShipCrewTracker: pa.last_event = event else: if pa is not None: del pa.v[pirate] - pa = PirateAboard(self._v, timestamp, event) + pa = PirateAboard(pirate, self._v, timestamp, event) self._pl[pirate] = pa self._v[pirate] = pa self._v['#lastaboard'] = timestamp self._refresh() return pa + def _trash_vessel(self, v): + for pn in v: + if pn.startswith('#'): continue + del self._pl[pn] + self._refresh() + + def expire_garbage(self, timestamp): + for (vn,v) in list(self._vl.iteritems()): + la = v['#lastaboard'] + if timestamp - la > opts.ship_reboard_clearout: + self._debug_line_disposition(timestamp,'', + 'stale reset '+vn) + self._trash_vessel(v) + del self._vl[vn] + def clear_vessel(self, timestamp): if self._v is not None: - for p in self._v: - if p.startswith('#'): continue - del self._pl[p] + self._trash_vessel(self._v) self._v = {'#lastaboard': timestamp} + self._vl[self._vessel] = self._v def _debug_line_disposition(self,timestamp,l,m): - debug('SCT %-13s %-30s %s' % (timestamp,m,l)) + debug('CLT %13s %-30s %s' % (timestamp,m,l)) def chatline(self,l): rm = lambda re: regexp.match(re,l) @@ -458,74 +493,181 @@ class ShipCrewTracker: timestamp = time.mktime(time_tuple) l = l[l.find(' ')+1:] - ob = lambda who, event: self._onboard_event( - timestamp, who, event) - oba = lambda m, did: ob( - m.group(1), '%s %s' % (did, m.group(2))) + def ob_x(who,event): + return self._onboard_event(timestamp, who, event) + def ob1(did): ob_x(m.group(1), did); return d(did) + def oba(did): return ob1('%s %s' % (did, m.group(2))) m = rm('Going aboard the (\\S.*\\S)\\.\\.\\.$') if m: + pn = self._myself.name self._vessel = m.group(1) dm = 'boarding' + try: self._v = self._vl[self._vessel] except KeyError: self._v = None; dm += ' new' + if self._v is not None: la = self._v['#lastaboard'] else: la = 0; dm += ' ?la' - if timestamp - la > 3600: + + if timestamp - la > opts.ship_reboard_clearout: self.clear_vessel(timestamp) dm += ' stale' - self._vl[self._vessel] = self._v - ob(self._myself.name, 'we boarded') + + ob_x(pn, 'we boarded') + self.expire_garbage(timestamp) return d(dm) if self._v is None: return d('no vessel') + m = rm('(\\w+) has come aboard\\.$') + if m: return ob1('boarded'); + m = rm('You have ordered (\\w+) to do some (\\S.*\\S)\\.$') if m: - pa = oba(m, 'ordered') - if m.group(2) == 'Gunning': + (who,what) = m.groups() + pa = ob_x(who,'ordered '+what) + if what == 'Gunning': pa.gunner = True return d('duty order') m = rm('(\\w+) abandoned a (\\S.*\\S) station\\.$') - if m: oba(m,'abandoned'); return d('abandoned') + if m: oba('stopped'); return d('stopped') + + def chat(what): + who = m.group(1) + try: pa = self._pl[who] + except KeyError: return d('chat mystery') + if pa.v is self._v: + pa.last_chat_time = timestamp + pa.last_chat_chan = what + self._refresh() + return d(what+' chat') + + m = rm('(\\w+) (?:issued an order|ordered everyone) "') + if m: return ob1('general order'); m = rm('(\\w+) says, "') - if m: ob(m.group(1), 'talked'); return d('talked') + if m: return chat('public') + + m = rm('(\\w+) tells ye, "') + if m: return chat('private') + + m = rm('(\\w+) flag officer chats, "') + if m: return chat('flag officer') + + m = rm('(\\w+) officer chats, "') + if m: return chat('officer') + + m = rm('Game over\\. Winners: ([A-Za-z, ]+)\\.$') + if m: + pl = m.group(1).split(', ') + if not self._myself.name in pl: + return d('lost boarding battle') + for pn in pl: + if ' ' in pn: continue + ob_x(pn,'won boarding battle') + return d('won boarding battle') + + m = rm('(\\w+) is eliminated\\!') + if m: return ob1('eliminated in fray'); m = rm('(\\w+) has left the vessel\.') if m: who = m.group(1) - ob(who, 'disembarked') + ob_x(who, 'disembarked') del self._v[who] del self._pl[who] return d('disembarked') return d('not matched') -def do_ship_aid(args, bu): - if len(args) != 1: bu('ship-aid takes only chat log filename') + def _str_vessel(self, vn, v): + s = ' vessel %s\n' % vn + s += ' '*20 + "%-*s %13s\n" % ( + max_pirate_namelen, '#lastaboard', + v['#lastaboard']) + for pn in sorted(v.keys()): + if pn.startswith('#'): continue + pa = v[pn] + assert pa.v == v + assert self._pl[pn] == pa + s += ' '*20 + "%s %-*s %13s %-30s %13s %s\n" % ( + (' ','G')[pa.gunner], + max_pirate_namelen, pn, + pa.last_time, pa.last_event, + pa.last_chat_time, pa.last_chat_chan) + return s + + def __str__(self): + s = '''