worker.go 9.5 KB

123456789101112131415161718192021222324252627282930313233343536373839404142434445464748495051525354555657585960616263646566676869707172737475767778798081828384858687888990919293949596979899100101102103104105106107108109110111112113114115116117118119120121122123124125126127128129130131132133134135136137138139140141142143144145146147148149150151152153154155156157158159160161162163164165166167168169170171172173174175176177178179180181182183184185186187188189190191192193194195196197198199200201202203204205206207208209210211212213214215216217218219220221222223224225226227228229230231232233234235236237238239240241242243244245246247248249250251252253254255256257258259260261262263264265266267268269270271272273274275276277278279280281282283284285286287288289290291292293294295296297298299300301302303304305306307308309310311312313314315316317318319320321322323324325326327328329330331332333334335336337338339340341342343344345346347348349350351352353354355356357358359360361362363364365366367368369370371372373374375376377378379380381382383384385386387388389390391392393394395396397398399400401402403404405406407408409410411412413414415416417418419420421422423
  1. package worker
  2. import (
  3. "fmt"
  4. "net/http"
  5. "sync"
  6. "time"
  7. "github.com/gin-gonic/gin"
  8. . "github.com/tuna/tunasync/internal"
  9. )
  10. var tunasyncWorker *Worker
  11. // A Worker is a instance of tunasync worker
  12. type Worker struct {
  13. L sync.Mutex
  14. cfg *Config
  15. jobs map[string]*mirrorJob
  16. managerChan chan jobMessage
  17. semaphore chan empty
  18. exit chan empty
  19. schedule *scheduleQueue
  20. httpEngine *gin.Engine
  21. httpClient *http.Client
  22. }
  23. // GetTUNASyncWorker returns a singalton worker
  24. func GetTUNASyncWorker(cfg *Config) *Worker {
  25. if tunasyncWorker != nil {
  26. return tunasyncWorker
  27. }
  28. w := &Worker{
  29. cfg: cfg,
  30. jobs: make(map[string]*mirrorJob),
  31. managerChan: make(chan jobMessage, 32),
  32. semaphore: make(chan empty, cfg.Global.Concurrent),
  33. exit: make(chan empty),
  34. schedule: newScheduleQueue(),
  35. }
  36. if cfg.Manager.CACert != "" {
  37. httpClient, err := CreateHTTPClient(cfg.Manager.CACert)
  38. if err != nil {
  39. logger.Errorf("Error initializing HTTP client: %s", err.Error())
  40. return nil
  41. }
  42. w.httpClient = httpClient
  43. }
  44. w.initJobs()
  45. w.makeHTTPServer()
  46. tunasyncWorker = w
  47. return w
  48. }
  49. func (w *Worker) initJobs() {
  50. for _, mirror := range w.cfg.Mirrors {
  51. // Create Provider
  52. provider := newMirrorProvider(mirror, w.cfg)
  53. w.jobs[provider.Name()] = newMirrorJob(provider)
  54. }
  55. }
  56. // ReloadMirrorConfig refresh the providers and jobs
  57. // from new mirror configs
  58. // TODO: deleted job should be removed from manager-side mirror list
  59. func (w *Worker) ReloadMirrorConfig(newMirrors []mirrorConfig) {
  60. w.L.Lock()
  61. defer w.L.Unlock()
  62. logger.Info("Reloading mirror configs")
  63. oldMirrors := w.cfg.Mirrors
  64. difference := diffMirrorConfig(oldMirrors, newMirrors)
  65. // first deal with deletion and modifications
  66. for _, op := range difference {
  67. if op.diffOp == diffAdd {
  68. continue
  69. }
  70. name := op.mirCfg.Name
  71. job, ok := w.jobs[name]
  72. if !ok {
  73. logger.Warningf("Job %s not found", name)
  74. continue
  75. }
  76. switch op.diffOp {
  77. case diffDelete:
  78. w.disableJob(job)
  79. delete(w.jobs, name)
  80. logger.Noticef("Deleted job %s", name)
  81. case diffModify:
  82. jobState := job.State()
  83. w.disableJob(job)
  84. // set new provider
  85. provider := newMirrorProvider(op.mirCfg, w.cfg)
  86. if err := job.SetProvider(provider); err != nil {
  87. logger.Errorf("Error setting job provider of %s: %s", name, err.Error())
  88. continue
  89. }
  90. // re-schedule job according to its previous state
  91. if jobState == stateDisabled {
  92. job.SetState(stateDisabled)
  93. } else if jobState == statePaused {
  94. job.SetState(statePaused)
  95. go job.Run(w.managerChan, w.semaphore)
  96. } else {
  97. job.SetState(stateNone)
  98. go job.Run(w.managerChan, w.semaphore)
  99. w.schedule.AddJob(time.Now(), job)
  100. }
  101. logger.Noticef("Reloaded job %s", name)
  102. }
  103. }
  104. // for added new jobs, just start new jobs
  105. for _, op := range difference {
  106. if op.diffOp != diffAdd {
  107. continue
  108. }
  109. provider := newMirrorProvider(op.mirCfg, w.cfg)
  110. job := newMirrorJob(provider)
  111. w.jobs[provider.Name()] = job
  112. job.SetState(stateNone)
  113. go job.Run(w.managerChan, w.semaphore)
  114. w.schedule.AddJob(time.Now(), job)
  115. logger.Noticef("New job %s", job.Name())
  116. }
  117. w.cfg.Mirrors = newMirrors
  118. }
  119. func (w *Worker) disableJob(job *mirrorJob) {
  120. w.schedule.Remove(job.Name())
  121. if job.State() != stateDisabled {
  122. job.ctrlChan <- jobDisable
  123. <-job.disabled
  124. }
  125. }
  126. // Ctrl server receives commands from the manager
  127. func (w *Worker) makeHTTPServer() {
  128. s := gin.New()
  129. s.Use(gin.Recovery())
  130. s.POST("/", func(c *gin.Context) {
  131. w.L.Lock()
  132. defer w.L.Unlock()
  133. var cmd WorkerCmd
  134. if err := c.BindJSON(&cmd); err != nil {
  135. c.JSON(http.StatusBadRequest, gin.H{"msg": "Invalid request"})
  136. return
  137. }
  138. job, ok := w.jobs[cmd.MirrorID]
  139. if !ok {
  140. c.JSON(http.StatusNotFound, gin.H{"msg": fmt.Sprintf("Mirror ``%s'' not found", cmd.MirrorID)})
  141. return
  142. }
  143. logger.Noticef("Received command: %v", cmd)
  144. // No matter what command, the existing job
  145. // schedule should be flushed
  146. w.schedule.Remove(job.Name())
  147. // if job disabled, start them first
  148. switch cmd.Cmd {
  149. case CmdStart, CmdRestart:
  150. if job.State() == stateDisabled {
  151. go job.Run(w.managerChan, w.semaphore)
  152. }
  153. }
  154. switch cmd.Cmd {
  155. case CmdStart:
  156. job.ctrlChan <- jobStart
  157. case CmdRestart:
  158. job.ctrlChan <- jobRestart
  159. case CmdStop:
  160. // if job is disabled, no goroutine would be there
  161. // receiving this signal
  162. if job.State() != stateDisabled {
  163. job.ctrlChan <- jobStop
  164. }
  165. case CmdDisable:
  166. w.disableJob(job)
  167. case CmdPing:
  168. // empty
  169. default:
  170. c.JSON(http.StatusNotAcceptable, gin.H{"msg": "Invalid Command"})
  171. return
  172. }
  173. c.JSON(http.StatusOK, gin.H{"msg": "OK"})
  174. })
  175. w.httpEngine = s
  176. }
  177. func (w *Worker) runHTTPServer() {
  178. addr := fmt.Sprintf("%s:%d", w.cfg.Server.Addr, w.cfg.Server.Port)
  179. httpServer := &http.Server{
  180. Addr: addr,
  181. Handler: w.httpEngine,
  182. ReadTimeout: 10 * time.Second,
  183. WriteTimeout: 10 * time.Second,
  184. }
  185. if w.cfg.Server.SSLCert == "" && w.cfg.Server.SSLKey == "" {
  186. if err := httpServer.ListenAndServe(); err != nil {
  187. panic(err)
  188. }
  189. } else {
  190. if err := httpServer.ListenAndServeTLS(w.cfg.Server.SSLCert, w.cfg.Server.SSLKey); err != nil {
  191. panic(err)
  192. }
  193. }
  194. }
  195. // Halt stops all jobs
  196. func (w *Worker) Halt() {
  197. w.L.Lock()
  198. logger.Notice("Stopping all the jobs")
  199. for _, job := range w.jobs {
  200. if job.State() != stateDisabled {
  201. job.ctrlChan <- jobHalt
  202. }
  203. }
  204. jobsDone.Wait()
  205. logger.Notice("All the jobs are stopped")
  206. w.L.Unlock()
  207. close(w.exit)
  208. }
  209. // Run runs worker forever
  210. func (w *Worker) Run() {
  211. w.registorWorker()
  212. go w.runHTTPServer()
  213. w.runSchedule()
  214. }
  215. func (w *Worker) runSchedule() {
  216. w.L.Lock()
  217. mirrorList := w.fetchJobStatus()
  218. unset := make(map[string]bool)
  219. for name := range w.jobs {
  220. unset[name] = true
  221. }
  222. // Fetch mirror list stored in the manager
  223. // put it on the scheduled time
  224. // if it's disabled, ignore it
  225. for _, m := range mirrorList {
  226. if job, ok := w.jobs[m.Name]; ok {
  227. delete(unset, m.Name)
  228. switch m.Status {
  229. case Disabled:
  230. job.SetState(stateDisabled)
  231. continue
  232. case Paused:
  233. job.SetState(statePaused)
  234. go job.Run(w.managerChan, w.semaphore)
  235. continue
  236. default:
  237. job.SetState(stateNone)
  238. go job.Run(w.managerChan, w.semaphore)
  239. stime := m.LastUpdate.Add(job.provider.Interval())
  240. logger.Debugf("Scheduling job %s @%s", job.Name(), stime.Format("2006-01-02 15:04:05"))
  241. w.schedule.AddJob(stime, job)
  242. }
  243. }
  244. }
  245. // some new jobs may be added
  246. // which does not exist in the
  247. // manager's mirror list
  248. for name := range unset {
  249. job := w.jobs[name]
  250. job.SetState(stateNone)
  251. go job.Run(w.managerChan, w.semaphore)
  252. w.schedule.AddJob(time.Now(), job)
  253. }
  254. w.L.Unlock()
  255. for {
  256. select {
  257. case jobMsg := <-w.managerChan:
  258. // got status update from job
  259. w.L.Lock()
  260. job, ok := w.jobs[jobMsg.name]
  261. w.L.Unlock()
  262. if !ok {
  263. logger.Warningf("Job %s not found", jobMsg.name)
  264. continue
  265. }
  266. if (job.State() != stateReady) && (job.State() != stateHalting) {
  267. logger.Infof("Job %s state is not ready, skip adding new schedule", jobMsg.name)
  268. continue
  269. }
  270. // syncing status is only meaningful when job
  271. // is running. If it's paused or disabled
  272. // a sync failure signal would be emitted
  273. // which needs to be ignored
  274. w.updateStatus(job, jobMsg)
  275. // only successful or the final failure msg
  276. // can trigger scheduling
  277. if jobMsg.schedule {
  278. schedTime := time.Now().Add(job.provider.Interval())
  279. logger.Noticef(
  280. "Next scheduled time for %s: %s",
  281. job.Name(),
  282. schedTime.Format("2006-01-02 15:04:05"),
  283. )
  284. w.schedule.AddJob(schedTime, job)
  285. }
  286. case <-time.Tick(5 * time.Second):
  287. // check schedule every 5 seconds
  288. if job := w.schedule.Pop(); job != nil {
  289. job.ctrlChan <- jobStart
  290. }
  291. case <-w.exit:
  292. // flush status update messages
  293. w.L.Lock()
  294. defer w.L.Unlock()
  295. for {
  296. select {
  297. case jobMsg := <-w.managerChan:
  298. logger.Debugf("status update from %s", jobMsg.name)
  299. job, ok := w.jobs[jobMsg.name]
  300. if !ok {
  301. continue
  302. }
  303. if jobMsg.status == Failed || jobMsg.status == Success {
  304. w.updateStatus(job, jobMsg)
  305. }
  306. default:
  307. return
  308. }
  309. }
  310. }
  311. }
  312. }
  313. // Name returns worker name
  314. func (w *Worker) Name() string {
  315. return w.cfg.Global.Name
  316. }
  317. // URL returns the url to http server of the worker
  318. func (w *Worker) URL() string {
  319. proto := "https"
  320. if w.cfg.Server.SSLCert == "" && w.cfg.Server.SSLKey == "" {
  321. proto = "http"
  322. }
  323. return fmt.Sprintf("%s://%s:%d/", proto, w.cfg.Server.Hostname, w.cfg.Server.Port)
  324. }
  325. func (w *Worker) registorWorker() {
  326. url := fmt.Sprintf(
  327. "%s/workers",
  328. w.cfg.Manager.APIBase,
  329. )
  330. msg := WorkerStatus{
  331. ID: w.Name(),
  332. URL: w.URL(),
  333. }
  334. if _, err := PostJSON(url, msg, w.httpClient); err != nil {
  335. logger.Errorf("Failed to register worker")
  336. }
  337. }
  338. func (w *Worker) updateStatus(job *mirrorJob, jobMsg jobMessage) {
  339. url := fmt.Sprintf(
  340. "%s/workers/%s/jobs/%s",
  341. w.cfg.Manager.APIBase,
  342. w.Name(),
  343. jobMsg.name,
  344. )
  345. p := job.provider
  346. smsg := MirrorStatus{
  347. Name: jobMsg.name,
  348. Worker: w.cfg.Global.Name,
  349. IsMaster: p.IsMaster(),
  350. Status: jobMsg.status,
  351. Upstream: p.Upstream(),
  352. Size: "unknown",
  353. ErrorMsg: jobMsg.msg,
  354. }
  355. if _, err := PostJSON(url, smsg, w.httpClient); err != nil {
  356. logger.Errorf("Failed to update mirror(%s) status: %s", jobMsg.name, err.Error())
  357. }
  358. }
  359. func (w *Worker) fetchJobStatus() []MirrorStatus {
  360. var mirrorList []MirrorStatus
  361. url := fmt.Sprintf(
  362. "%s/workers/%s/jobs",
  363. w.cfg.Manager.APIBase,
  364. w.Name(),
  365. )
  366. if _, err := GetJSON(url, &mirrorList, w.httpClient); err != nil {
  367. logger.Errorf("Failed to fetch job status: %s", err.Error())
  368. }
  369. return mirrorList
  370. }