2 * SPDX-License-Identifier: BSD-2-Clause-FreeBSD
4 * Copyright (c) 2011 NetApp, Inc.
7 * Redistribution and use in source and binary forms, with or without
8 * modification, are permitted provided that the following conditions
10 * 1. Redistributions of source code must retain the above copyright
11 * notice, this list of conditions and the following disclaimer.
12 * 2. Redistributions in binary form must reproduce the above copyright
13 * notice, this list of conditions and the following disclaimer in the
14 * documentation and/or other materials provided with the distribution.
16 * THIS SOFTWARE IS PROVIDED BY NETAPP, INC ``AS IS'' AND
17 * ANY EXPRESS OR IMPLIED WARRANTIES, INCLUDING, BUT NOT LIMITED TO, THE
18 * IMPLIED WARRANTIES OF MERCHANTABILITY AND FITNESS FOR A PARTICULAR PURPOSE
19 * ARE DISCLAIMED. IN NO EVENT SHALL NETAPP, INC OR CONTRIBUTORS BE LIABLE
20 * FOR ANY DIRECT, INDIRECT, INCIDENTAL, SPECIAL, EXEMPLARY, OR CONSEQUENTIAL
21 * DAMAGES (INCLUDING, BUT NOT LIMITED TO, PROCUREMENT OF SUBSTITUTE GOODS
22 * OR SERVICES; LOSS OF USE, DATA, OR PROFITS; OR BUSINESS INTERRUPTION)
23 * HOWEVER CAUSED AND ON ANY THEORY OF LIABILITY, WHETHER IN CONTRACT, STRICT
24 * LIABILITY, OR TORT (INCLUDING NEGLIGENCE OR OTHERWISE) ARISING IN ANY WAY
25 * OUT OF THE USE OF THIS SOFTWARE, EVEN IF ADVISED OF THE POSSIBILITY OF
32 * Micro event library for FreeBSD, designed for a single i/o thread
33 * using kqueue, and having events be persistent by default.
36 #include <sys/cdefs.h>
37 __FBSDID("$FreeBSD$");
40 #ifndef WITHOUT_CAPSICUM
41 #include <capsicum_helpers.h>
52 #include <sys/types.h>
53 #ifndef WITHOUT_CAPSICUM
54 #include <sys/capsicum.h>
56 #include <sys/event.h>
60 #include <pthread_np.h>
66 static pthread_t mevent_tid;
67 static pthread_once_t mevent_once = PTHREAD_ONCE_INIT;
68 static int mevent_timid = 43;
69 static int mevent_pipefd[2];
71 static pthread_mutex_t mevent_lmutex = PTHREAD_MUTEX_INITIALIZER;
74 void (*me_func)(int, enum ev_type, void *);
75 #define me_msecs me_fd
81 int me_state; /* Desired kevent flags. */
84 LIST_ENTRY(mevent) me_list;
87 static LIST_HEAD(listhead, mevent) global_head, change_head;
92 pthread_mutex_lock(&mevent_lmutex);
98 pthread_mutex_unlock(&mevent_lmutex);
102 mevent_pipe_read(int fd, enum ev_type type, void *param)
104 char buf[MEVENT_MAX];
108 * Drain the pipe read side. The fd is non-blocking so this is
112 status = read(fd, buf, sizeof(buf));
113 } while (status == MEVENT_MAX);
122 * If calling from outside the i/o thread, write a byte on the
123 * pipe to force the i/o thread to exit the blocking kevent call.
125 if (mevent_pipefd[1] != 0 && pthread_self() != mevent_tid) {
126 write(mevent_pipefd[1], &c, 1);
133 #ifndef WITHOUT_CAPSICUM
140 #ifndef WITHOUT_CAPSICUM
141 cap_rights_init(&rights, CAP_KQUEUE);
142 if (caph_rights_limit(mfd, &rights) == -1)
143 errx(EX_OSERR, "Unable to apply rights for sandbox");
146 LIST_INIT(&change_head);
147 LIST_INIT(&global_head);
151 mevent_kq_filter(struct mevent *mevp)
157 if (mevp->me_type == EVF_READ)
158 retval = EVFILT_READ;
160 if (mevp->me_type == EVF_WRITE)
161 retval = EVFILT_WRITE;
163 if (mevp->me_type == EVF_TIMER)
164 retval = EVFILT_TIMER;
166 if (mevp->me_type == EVF_SIGNAL)
167 retval = EVFILT_SIGNAL;
169 if (mevp->me_type == EVF_VNODE)
170 retval = EVFILT_VNODE;
176 mevent_kq_flags(struct mevent *mevp)
180 retval = mevp->me_state;
182 if (mevp->me_type == EVF_VNODE)
189 mevent_kq_fflags(struct mevent *mevp)
195 switch (mevp->me_type) {
197 if ((mevp->me_fflags & EVFF_ATTRIB) != 0)
198 retval |= NOTE_ATTRIB;
211 mevent_populate(struct mevent *mevp, struct kevent *kev)
213 if (mevp->me_type == EVF_TIMER) {
214 kev->ident = mevp->me_timid;
215 kev->data = mevp->me_msecs;
217 kev->ident = mevp->me_fd;
220 kev->filter = mevent_kq_filter(mevp);
221 kev->flags = mevent_kq_flags(mevp);
222 kev->fflags = mevent_kq_fflags(mevp);
227 mevent_build(struct kevent *kev)
229 struct mevent *mevp, *tmpp;
236 LIST_FOREACH_SAFE(mevp, &change_head, me_list, tmpp) {
237 if (mevp->me_closefd) {
239 * A close of the file descriptor will remove the
244 assert((mevp->me_state & EV_ADD) == 0);
245 mevent_populate(mevp, &kev[i]);
250 LIST_REMOVE(mevp, me_list);
252 if (mevp->me_state & EV_DELETE) {
255 LIST_INSERT_HEAD(&global_head, mevp, me_list);
258 assert(i < MEVENT_MAX);
267 mevent_handle(struct kevent *kev, int numev)
272 for (i = 0; i < numev; i++) {
275 /* XXX check for EV_ERROR ? */
277 (*mevp->me_func)(mevp->me_fd, mevp->me_type, mevp->me_param);
281 static struct mevent *
282 mevent_add_state(int tfd, enum ev_type type,
283 void (*func)(int, enum ev_type, void *), void *param,
284 int state, int fflags)
287 struct mevent *lp, *mevp;
290 if (tfd < 0 || func == NULL) {
296 pthread_once(&mevent_once, mevent_init);
301 * Verify that the fd/type tuple is not present in any list
303 LIST_FOREACH(lp, &global_head, me_list) {
304 if (type != EVF_TIMER && lp->me_fd == tfd &&
305 lp->me_type == type) {
310 LIST_FOREACH(lp, &change_head, me_list) {
311 if (type != EVF_TIMER && lp->me_fd == tfd &&
312 lp->me_type == type) {
318 * Allocate an entry and populate it.
320 mevp = calloc(1, sizeof(struct mevent));
325 if (type == EVF_TIMER) {
326 mevp->me_msecs = tfd;
327 mevp->me_timid = mevent_timid++;
330 mevp->me_type = type;
331 mevp->me_func = func;
332 mevp->me_param = param;
333 mevp->me_state = state;
334 mevp->me_fflags = fflags;
337 * Try to add the event. If this fails, report the failure to
340 mevent_populate(mevp, &kev);
341 ret = kevent(mfd, &kev, 1, NULL, 0, NULL);
348 mevp->me_state &= ~EV_ADD;
349 LIST_INSERT_HEAD(&global_head, mevp, me_list);
358 mevent_add(int tfd, enum ev_type type,
359 void (*func)(int, enum ev_type, void *), void *param)
362 return (mevent_add_state(tfd, type, func, param, EV_ADD, 0));
366 mevent_add_flags(int tfd, enum ev_type type, int fflags,
367 void (*func)(int, enum ev_type, void *), void *param)
370 return (mevent_add_state(tfd, type, func, param, EV_ADD, fflags));
374 mevent_add_disabled(int tfd, enum ev_type type,
375 void (*func)(int, enum ev_type, void *), void *param)
378 return (mevent_add_state(tfd, type, func, param, EV_ADD | EV_DISABLE, 0));
382 mevent_update(struct mevent *evp, bool enable)
389 * It's not possible to enable/disable a deleted event
391 assert((evp->me_state & EV_DELETE) == 0);
393 newstate = evp->me_state;
395 newstate |= EV_ENABLE;
396 newstate &= ~EV_DISABLE;
398 newstate |= EV_DISABLE;
399 newstate &= ~EV_ENABLE;
403 * No update needed if state isn't changing
405 if (evp->me_state != newstate) {
406 evp->me_state = newstate;
409 * Place the entry onto the changed list if not
412 if (evp->me_cq == 0) {
414 LIST_REMOVE(evp, me_list);
415 LIST_INSERT_HEAD(&change_head, evp, me_list);
426 mevent_enable(struct mevent *evp)
429 return (mevent_update(evp, true));
433 mevent_disable(struct mevent *evp)
436 return (mevent_update(evp, false));
440 mevent_delete_event(struct mevent *evp, int closefd)
445 * Place the entry onto the changed list if not already there, and
446 * mark as to be deleted.
448 if (evp->me_cq == 0) {
450 LIST_REMOVE(evp, me_list);
451 LIST_INSERT_HEAD(&change_head, evp, me_list);
454 evp->me_state = EV_DELETE;
465 mevent_delete(struct mevent *evp)
468 return (mevent_delete_event(evp, 0));
472 mevent_delete_close(struct mevent *evp)
475 return (mevent_delete_event(evp, 1));
479 mevent_set_name(void)
482 pthread_set_name_np(mevent_tid, "mevent");
486 mevent_dispatch(void)
488 struct kevent changelist[MEVENT_MAX];
489 struct kevent eventlist[MEVENT_MAX];
490 struct mevent *pipev;
493 #ifndef WITHOUT_CAPSICUM
497 mevent_tid = pthread_self();
500 pthread_once(&mevent_once, mevent_init);
503 * Open the pipe that will be used for other threads to force
504 * the blocking kqueue call to exit by writing to it. Set the
505 * descriptor to non-blocking.
507 ret = pipe(mevent_pipefd);
513 #ifndef WITHOUT_CAPSICUM
514 cap_rights_init(&rights, CAP_EVENT, CAP_READ, CAP_WRITE);
515 if (caph_rights_limit(mevent_pipefd[0], &rights) == -1)
516 errx(EX_OSERR, "Unable to apply rights for sandbox");
517 if (caph_rights_limit(mevent_pipefd[1], &rights) == -1)
518 errx(EX_OSERR, "Unable to apply rights for sandbox");
522 * Add internal event handler for the pipe write fd
524 pipev = mevent_add(mevent_pipefd[0], EVF_READ, mevent_pipe_read, NULL);
525 assert(pipev != NULL);
529 * Build changelist if required.
530 * XXX the changelist can be put into the blocking call
531 * to eliminate the extra syscall. Currently better for
534 numev = mevent_build(changelist);
536 ret = kevent(mfd, changelist, numev, NULL, 0, NULL);
538 perror("Error return from kevent change");
543 * Block awaiting events
545 ret = kevent(mfd, NULL, 0, eventlist, MEVENT_MAX, NULL);
546 if (ret == -1 && errno != EINTR) {
547 perror("Error return from kevent monitor");
551 * Handle reported events
553 mevent_handle(eventlist, ret);