summaryrefslogtreecommitdiff
path: root/de-sgm
blob: 452edfe6a64a747502bf0000de30cf115ac1fd28 (plain)
1
2
3
4
5
6
7
8
9
10
#!/bin/sh

egrep -v "^[[:space:]]*(<\?xml.*\?>|</?(mteval|doc|srcset|refset|translator|reviewer)[^>]*>)[[:space:]]*$" \
  | egrep -v "^[[:space:]]*<(url|description|keywords|talkid|title|translator|reviewer)[^>]*>.*</(url|description|keywords|talkid|title|translator|reviewer)>[[:space:]]*$" \
  | sed "s|<seg[^>]*>\s*||" \
  | sed "s|\s*</seg>\s*$||" \
  | egrep -v "^[[:space:]]*<p>[[:space:]]*$|^[[:space:]]*</p>[[:space:]]*$" \
  | sed "s|<speaker>\s*||" \
  | sed "s|\s*</speaker>\s*$||"