2004-07-01 18:49:54 +04:00
|
|
|
/*
|
2005-11-05 22:57:48 +03:00
|
|
|
* Copyright (c) 2004-2005 The Trustees of Indiana University and Indiana
|
|
|
|
* University Research and Technology
|
|
|
|
* Corporation. All rights reserved.
|
|
|
|
* Copyright (c) 2004-2005 The University of Tennessee and The University
|
|
|
|
* of Tennessee Research Foundation. All rights
|
|
|
|
* reserved.
|
2004-11-28 23:09:25 +03:00
|
|
|
* Copyright (c) 2004-2005 High Performance Computing Center Stuttgart,
|
|
|
|
* University of Stuttgart. All rights reserved.
|
2005-03-24 15:43:37 +03:00
|
|
|
* Copyright (c) 2004-2005 The Regents of the University of California.
|
|
|
|
* All rights reserved.
|
2004-11-22 04:38:40 +03:00
|
|
|
* $COPYRIGHT$
|
|
|
|
*
|
|
|
|
* Additional copyrights may follow
|
|
|
|
*
|
2004-07-01 18:49:54 +04:00
|
|
|
* $HEADER$
|
|
|
|
*/
|
|
|
|
/** @file:
|
|
|
|
*
|
|
|
|
* Defines the functions for the tcp module.
|
|
|
|
*/
|
|
|
|
|
|
|
|
#ifndef _MCA_OOB_TCP_H_
|
|
|
|
#define _MCA_OOB_TCP_H_
|
|
|
|
|
|
|
|
#include "mca/oob/oob.h"
|
|
|
|
#include "mca/oob/base/base.h"
|
2004-08-05 03:42:51 +04:00
|
|
|
#include "mca/base/base.h"
|
2005-03-14 23:57:21 +03:00
|
|
|
#include "mca/ns/ns_types.h"
|
2005-07-02 20:46:27 +04:00
|
|
|
#include "opal/class/opal_free_list.h"
|
2005-07-03 20:52:32 +04:00
|
|
|
#include "class/opal_hash_table.h"
|
2005-07-04 03:09:55 +04:00
|
|
|
#include "opal/event/event.h"
|
2005-07-04 02:45:48 +04:00
|
|
|
#include "opal/threads/mutex.h"
|
|
|
|
#include "opal/threads/condition.h"
|
2004-07-01 18:49:54 +04:00
|
|
|
#include "mca/oob/tcp/oob_tcp_peer.h"
|
2004-07-13 02:46:57 +04:00
|
|
|
#include "mca/oob/tcp/oob_tcp_msg.h"
|
2004-07-01 18:49:54 +04:00
|
|
|
|
|
|
|
|
|
|
|
#if defined(c_plusplus) || defined(__cplusplus)
|
|
|
|
extern "C" {
|
|
|
|
#endif
|
|
|
|
|
2004-09-02 03:07:40 +04:00
|
|
|
|
|
|
|
|
2004-07-01 18:49:54 +04:00
|
|
|
/*
|
2004-08-19 23:34:37 +04:00
|
|
|
* standard component functions
|
2004-07-01 18:49:54 +04:00
|
|
|
*/
|
2004-08-19 23:34:37 +04:00
|
|
|
int mca_oob_tcp_component_open(void);
|
|
|
|
int mca_oob_tcp_component_close(void);
|
2005-03-14 23:57:21 +03:00
|
|
|
mca_oob_t* mca_oob_tcp_component_init(int* priority);
|
2004-08-19 23:34:37 +04:00
|
|
|
|
|
|
|
/**
|
|
|
|
* Hook function to allow the selected oob components
|
|
|
|
* to register their contact info with the registry
|
|
|
|
*/
|
|
|
|
|
|
|
|
int mca_oob_tcp_init(void);
|
|
|
|
|
|
|
|
/**
|
|
|
|
* Cleanup resources during shutdown.
|
|
|
|
*/
|
|
|
|
int mca_oob_tcp_fini(void);
|
2004-07-01 18:49:54 +04:00
|
|
|
|
2005-07-21 21:45:09 +04:00
|
|
|
#if SIZEOF_SIZE_T == 8
|
|
|
|
/*
|
|
|
|
* Convert a 64 bit value to network byte order.
|
|
|
|
*/
|
|
|
|
static inline uint64_t hton64(uint64_t val)
|
|
|
|
{
|
|
|
|
union { uint64_t ll;
|
|
|
|
uint32_t l[2];
|
|
|
|
} w, r;
|
|
|
|
|
|
|
|
/* platform already in network byte order? */
|
|
|
|
if(htonl(1) == 1L)
|
|
|
|
return val;
|
|
|
|
w.ll = val;
|
|
|
|
r.l[0] = htonl(w.l[1]);
|
|
|
|
r.l[1] = htonl(w.l[0]);
|
|
|
|
return r.ll;
|
|
|
|
}
|
|
|
|
|
|
|
|
/*
|
|
|
|
* Convert a 64 bit value from network to host byte order.
|
|
|
|
*/
|
|
|
|
|
|
|
|
static inline uint64_t ntoh64(uint64_t val)
|
|
|
|
{
|
|
|
|
union { uint64_t ll;
|
|
|
|
uint32_t l[2];
|
|
|
|
} w, r;
|
|
|
|
|
|
|
|
/* platform already in network byte order? */
|
|
|
|
if(htonl(1) == 1L)
|
|
|
|
return val;
|
|
|
|
w.ll = val;
|
|
|
|
r.l[0] = ntohl(w.l[1]);
|
|
|
|
r.l[1] = ntohl(w.l[0]);
|
|
|
|
return r.ll;
|
|
|
|
}
|
|
|
|
|
|
|
|
/**
|
|
|
|
* Convert process name from host to network byte order.
|
|
|
|
*
|
|
|
|
* @param name
|
|
|
|
*/
|
|
|
|
#define OMPI_PROCESS_NAME_HTON(n) \
|
|
|
|
n.cellid = hton64(n.cellid); \
|
|
|
|
n.jobid = hton64(n.jobid); \
|
|
|
|
n.vpid = hton64(n.vpid);
|
|
|
|
|
2004-08-10 03:07:53 +04:00
|
|
|
/**
|
2004-08-03 01:24:00 +04:00
|
|
|
* Convert process name from network to host byte order.
|
|
|
|
*
|
|
|
|
* @param name
|
|
|
|
*/
|
|
|
|
#define OMPI_PROCESS_NAME_NTOH(n) \
|
2005-07-21 21:45:09 +04:00
|
|
|
n.cellid = ntoh64(n.cellid); \
|
|
|
|
n.jobid = ntoh64(n.jobid); \
|
|
|
|
n.vpid = ntoh64(n.vpid);
|
|
|
|
|
|
|
|
#else
|
2004-07-01 18:49:54 +04:00
|
|
|
|
2004-08-10 03:07:53 +04:00
|
|
|
/**
|
2004-08-03 01:24:00 +04:00
|
|
|
* Convert process name from host to network byte order.
|
|
|
|
*
|
|
|
|
* @param name
|
|
|
|
*/
|
|
|
|
#define OMPI_PROCESS_NAME_HTON(n) \
|
|
|
|
n.cellid = htonl(n.cellid); \
|
|
|
|
n.jobid = htonl(n.jobid); \
|
|
|
|
n.vpid = htonl(n.vpid);
|
|
|
|
|
2005-07-21 21:45:09 +04:00
|
|
|
/**
|
|
|
|
* Convert process name from network to host byte order.
|
|
|
|
*
|
|
|
|
* @param name
|
|
|
|
*/
|
2005-07-22 00:18:39 +04:00
|
|
|
#define OMPI_PROCESS_NAME_NTOH(n) \
|
2005-07-21 21:45:09 +04:00
|
|
|
n.cellid = ntohl(n.cellid); \
|
|
|
|
n.jobid = ntohl(n.jobid); \
|
|
|
|
n.vpid = ntohl(n.vpid);
|
2005-07-22 00:18:39 +04:00
|
|
|
|
2005-07-21 21:45:09 +04:00
|
|
|
#endif
|
2004-08-03 01:24:00 +04:00
|
|
|
|
|
|
|
/**
|
|
|
|
* Compare two process names for equality.
|
|
|
|
*
|
|
|
|
* @param n1 Process name 1.
|
|
|
|
* @param n2 Process name 2.
|
|
|
|
* @return (-1 for n1<n2 0 for equality, 1 for n1>n2)
|
|
|
|
*
|
|
|
|
* Note that the definition of < or > is somewhat arbitrary -
|
|
|
|
* just needs to be consistently applied to maintain an ordering
|
|
|
|
* when process names are used as indices.
|
|
|
|
*/
|
2005-03-14 23:57:21 +03:00
|
|
|
int mca_oob_tcp_process_name_compare(const orte_process_name_t* n1, const orte_process_name_t* n2);
|
2005-09-15 21:13:13 +04:00
|
|
|
|
|
|
|
/**
|
2004-08-16 23:39:54 +04:00
|
|
|
* Obtain contact information for this host (e.g. <ipaddress>:<port>)
|
|
|
|
*/
|
|
|
|
|
|
|
|
char* mca_oob_tcp_get_addr(void);
|
|
|
|
|
|
|
|
/**
|
Not as bad as this all may look. Tim and I made a significant change to the way we handle the startup of the oob, the seed, etc. We have made it backwards-compatible so that mpirun2 and singleton operations remain working. We had to adjust the name server and gpr as well, plus the process_info structure.
This also includes a checkpoint update to openmpi.c and ompid.c. I have re-enabled the ompid compile.
This latter raises an important point. The trunk compiles the programs like ompid just fine under Linux. It also does just fine for OSX under the dynamic libraries. However, we are seeing errors when compiling under OSX for the static case - the linker seems to have trouble resolving some variable names, even though linker diagnostics show the variables as being defined. Thus, a warning to Mac users that you may have to locally turn things off if you are trying to do static compiles. We ask, however, that you don't commit those changes that turn things off for everyone else - instead, let's try to figure out why the static compile is having a problem, and let everyone else continue to work.
Thanks
Ralph
This commit was SVN r2534.
2004-09-08 07:59:06 +04:00
|
|
|
* Setup cached addresses for the peers.
|
2004-08-16 23:39:54 +04:00
|
|
|
*/
|
|
|
|
|
2005-03-14 23:57:21 +03:00
|
|
|
int mca_oob_tcp_set_addr(const orte_process_name_t*, const char*);
|
2004-08-16 23:39:54 +04:00
|
|
|
|
2004-09-08 21:02:24 +04:00
|
|
|
/**
|
|
|
|
* A routine to ping a given process name to determine if it is reachable.
|
|
|
|
*
|
|
|
|
* @param name The peer name.
|
|
|
|
* @param tv The length of time to wait on a connection/response.
|
|
|
|
*
|
|
|
|
* Note that this routine blocks up to the specified timeout waiting for a
|
|
|
|
* connection / response from the specified peer. If the peer is unavailable
|
|
|
|
* an error status is returned.
|
|
|
|
*/
|
|
|
|
|
2005-05-05 20:31:40 +04:00
|
|
|
int mca_oob_tcp_ping(const orte_process_name_t*, const char* uri, const struct timeval* tv);
|
2004-09-08 21:02:24 +04:00
|
|
|
|
2004-07-01 18:49:54 +04:00
|
|
|
/**
|
2004-07-14 01:03:03 +04:00
|
|
|
* Similiar to unix writev(2).
|
2004-07-01 18:49:54 +04:00
|
|
|
*
|
|
|
|
* @param peer (IN) Opaque name of peer process.
|
|
|
|
* @param msg (IN) Array of iovecs describing user buffers and lengths.
|
|
|
|
* @param count (IN) Number of elements in iovec array.
|
2004-08-03 01:24:00 +04:00
|
|
|
* @param tag (IN) User defined tag for matching send/recv.
|
2004-07-01 18:49:54 +04:00
|
|
|
* @param flags (IN) Currently unused.
|
|
|
|
* @return OMPI error code (<0) on error number of bytes actually sent.
|
|
|
|
*/
|
|
|
|
|
2004-08-03 01:24:00 +04:00
|
|
|
int mca_oob_tcp_send(
|
2005-03-14 23:57:21 +03:00
|
|
|
orte_process_name_t* peer,
|
2004-08-13 02:41:42 +04:00
|
|
|
struct iovec *msg,
|
2004-08-03 01:24:00 +04:00
|
|
|
int count,
|
|
|
|
int tag,
|
|
|
|
int flags);
|
2004-07-01 18:49:54 +04:00
|
|
|
|
|
|
|
/**
|
2004-07-14 01:03:03 +04:00
|
|
|
* Similiar to unix readv(2)
|
2004-07-01 18:49:54 +04:00
|
|
|
*
|
2004-08-05 03:42:51 +04:00
|
|
|
* @param peer (IN) Opaque name of peer process or MCA_OOB_NAME_ANY for wildcard receive.
|
2004-07-01 18:49:54 +04:00
|
|
|
* @param msg (IN) Array of iovecs describing user buffers and lengths.
|
|
|
|
* @param count (IN) Number of elements in iovec array.
|
2004-08-03 01:24:00 +04:00
|
|
|
* @param tag (IN) User defined tag for matching send/recv.
|
2004-07-15 23:08:54 +04:00
|
|
|
* @param flags (IN) May be MCA_OOB_PEEK to return up to the number of bytes provided in the
|
2004-07-01 18:49:54 +04:00
|
|
|
* iovec array without removing the message from the queue.
|
|
|
|
* @return OMPI error code (<0) on error or number of bytes actually received.
|
|
|
|
*/
|
|
|
|
|
2004-08-03 01:24:00 +04:00
|
|
|
int mca_oob_tcp_recv(
|
2005-03-14 23:57:21 +03:00
|
|
|
orte_process_name_t* peer,
|
2004-08-13 02:41:42 +04:00
|
|
|
struct iovec * msg,
|
2004-08-03 01:24:00 +04:00
|
|
|
int count,
|
2005-03-14 23:57:21 +03:00
|
|
|
int tag,
|
2004-08-03 01:24:00 +04:00
|
|
|
int flags);
|
2004-07-01 18:49:54 +04:00
|
|
|
|
|
|
|
|
|
|
|
/*
|
|
|
|
* Non-blocking versions of send/recv.
|
|
|
|
*/
|
|
|
|
|
|
|
|
/**
|
|
|
|
* Non-blocking version of mca_oob_send().
|
|
|
|
*
|
|
|
|
* @param peer (IN) Opaque name of peer process.
|
|
|
|
* @param msg (IN) Array of iovecs describing user buffers and lengths.
|
|
|
|
* @param count (IN) Number of elements in iovec array.
|
2004-08-03 01:24:00 +04:00
|
|
|
* @param tag (IN) User defined tag for matching send/recv.
|
2004-07-01 18:49:54 +04:00
|
|
|
* @param flags (IN) Currently unused.
|
|
|
|
* @param cbfunc (IN) Callback function on send completion.
|
|
|
|
* @param cbdata (IN) User data that is passed to callback function.
|
|
|
|
* @return OMPI error code (<0) on error number of bytes actually sent.
|
|
|
|
*
|
|
|
|
*/
|
|
|
|
|
2004-08-03 01:24:00 +04:00
|
|
|
int mca_oob_tcp_send_nb(
|
2005-03-14 23:57:21 +03:00
|
|
|
orte_process_name_t* peer,
|
2004-08-13 02:41:42 +04:00
|
|
|
struct iovec* msg,
|
2004-08-03 01:24:00 +04:00
|
|
|
int count,
|
|
|
|
int tag,
|
|
|
|
int flags,
|
|
|
|
mca_oob_callback_fn_t cbfunc,
|
|
|
|
void* cbdata);
|
2004-07-01 18:49:54 +04:00
|
|
|
|
|
|
|
/**
|
|
|
|
* Non-blocking version of mca_oob_recv().
|
|
|
|
*
|
2004-08-05 03:42:51 +04:00
|
|
|
* @param peer (IN) Opaque name of peer process or MCA_OOB_NAME_ANY for wildcard receive.
|
2004-07-01 18:49:54 +04:00
|
|
|
* @param msg (IN) Array of iovecs describing user buffers and lengths.
|
|
|
|
* @param count (IN) Number of elements in iovec array.
|
2004-08-03 01:24:00 +04:00
|
|
|
* @param tag (IN) User defined tag for matching send/recv.
|
2004-07-15 23:08:54 +04:00
|
|
|
* @param flags (IN) May be MCA_OOB_PEEK to return up to size bytes of msg w/out removing it from the queue,
|
2004-07-01 18:49:54 +04:00
|
|
|
* @param cbfunc (IN) Callback function on recv completion.
|
|
|
|
* @param cbdata (IN) User data that is passed to callback function.
|
|
|
|
* @return OMPI error code (<0) on error or number of bytes actually received.
|
|
|
|
*/
|
|
|
|
|
2004-08-03 01:24:00 +04:00
|
|
|
int mca_oob_tcp_recv_nb(
|
2005-03-14 23:57:21 +03:00
|
|
|
orte_process_name_t* peer,
|
2004-08-13 02:41:42 +04:00
|
|
|
struct iovec* msg,
|
2004-08-03 01:24:00 +04:00
|
|
|
int count,
|
|
|
|
int tag,
|
|
|
|
int flags,
|
|
|
|
mca_oob_callback_fn_t cbfunc,
|
|
|
|
void* cbdata);
|
2004-07-01 18:49:54 +04:00
|
|
|
|
2004-09-30 19:09:29 +04:00
|
|
|
/**
|
|
|
|
* Cancel non-blocking receive.
|
|
|
|
*
|
|
|
|
* @param peer (IN) Opaque name of peer process or MCA_OOB_NAME_ANY for wildcard receive.
|
|
|
|
* @param tag (IN) User defined tag for matching send/recv.
|
|
|
|
* @return OMPI error code (<0) on error or number of bytes actually received.
|
|
|
|
*/
|
|
|
|
|
|
|
|
int mca_oob_tcp_recv_cancel(
|
2005-03-14 23:57:21 +03:00
|
|
|
orte_process_name_t* peer,
|
2004-09-30 19:09:29 +04:00
|
|
|
int tag);
|
|
|
|
|
2004-09-02 03:07:40 +04:00
|
|
|
/**
|
|
|
|
* Attempt to map a peer name to its corresponding address.
|
|
|
|
*/
|
|
|
|
|
|
|
|
int mca_oob_tcp_resolve(mca_oob_tcp_peer_t*);
|
|
|
|
|
2004-08-19 23:34:37 +04:00
|
|
|
/**
|
|
|
|
* Parse a URI string into an IP address and port number.
|
|
|
|
*/
|
|
|
|
int mca_oob_tcp_parse_uri(
|
|
|
|
const char* uri,
|
|
|
|
struct sockaddr_in* inaddr
|
|
|
|
);
|
|
|
|
|
2004-11-20 22:12:43 +03:00
|
|
|
/**
|
|
|
|
* Callback from registry on change to subscribed segments
|
|
|
|
*/
|
|
|
|
void mca_oob_tcp_registry_callback(
|
2005-03-14 23:57:21 +03:00
|
|
|
orte_gpr_notify_data_t* data,
|
2004-11-20 22:12:43 +03:00
|
|
|
void* cbdata);
|
|
|
|
|
2005-10-31 19:21:11 +03:00
|
|
|
/**
|
|
|
|
* Setup socket options
|
|
|
|
*/
|
|
|
|
|
|
|
|
void mca_oob_tcp_set_socket_options(int sd);
|
2004-07-01 18:49:54 +04:00
|
|
|
|
2004-07-13 02:46:57 +04:00
|
|
|
/**
|
|
|
|
* OOB TCP Component
|
|
|
|
*/
|
|
|
|
struct mca_oob_tcp_component_t {
|
2004-07-15 17:51:40 +04:00
|
|
|
mca_oob_base_component_1_0_0_t super; /**< base OOB component */
|
2005-03-19 02:40:08 +03:00
|
|
|
char* tcp_include; /**< list of ip interfaces to include */
|
|
|
|
char* tcp_exclude; /**< list of ip interfaces to exclude */
|
2004-08-16 23:39:54 +04:00
|
|
|
int tcp_listen_sd; /**< listen socket for incoming connection requests */
|
|
|
|
unsigned short tcp_listen_port; /**< listen port */
|
2005-07-03 20:22:16 +04:00
|
|
|
opal_list_t tcp_subscriptions; /**< list of registry subscriptions */
|
|
|
|
opal_list_t tcp_peer_list; /**< list of peers sorted in mru order */
|
2005-07-03 20:52:32 +04:00
|
|
|
opal_hash_table_t tcp_peers; /**< peers sorted by name */
|
|
|
|
opal_hash_table_t tcp_peer_names; /**< cache of peer contact info sorted by name */
|
2005-07-02 20:46:27 +04:00
|
|
|
opal_free_list_t tcp_peer_free; /**< free list of peers */
|
Not as bad as this all may look. Tim and I made a significant change to the way we handle the startup of the oob, the seed, etc. We have made it backwards-compatible so that mpirun2 and singleton operations remain working. We had to adjust the name server and gpr as well, plus the process_info structure.
This also includes a checkpoint update to openmpi.c and ompid.c. I have re-enabled the ompid compile.
This latter raises an important point. The trunk compiles the programs like ompid just fine under Linux. It also does just fine for OSX under the dynamic libraries. However, we are seeing errors when compiling under OSX for the static case - the linker seems to have trouble resolving some variable names, even though linker diagnostics show the variables as being defined. Thus, a warning to Mac users that you may have to locally turn things off if you are trying to do static compiles. We ask, however, that you don't commit those changes that turn things off for everyone else - instead, let's try to figure out why the static compile is having a problem, and let everyone else continue to work.
Thanks
Ralph
This commit was SVN r2534.
2004-09-08 07:59:06 +04:00
|
|
|
int tcp_peer_limit; /**< max size of tcp peer cache */
|
2004-08-16 23:39:54 +04:00
|
|
|
int tcp_peer_retries; /**< max number of retries before declaring peer gone */
|
2005-10-31 19:21:11 +03:00
|
|
|
int tcp_sndbuf; /**< socket send buffer size */
|
|
|
|
int tcp_rcvbuf; /**< socket recv buffer size */
|
2005-07-02 20:46:27 +04:00
|
|
|
opal_free_list_t tcp_msgs; /**< free list of messages */
|
2005-07-04 03:09:55 +04:00
|
|
|
opal_event_t tcp_send_event; /**< event structure for sends */
|
|
|
|
opal_event_t tcp_recv_event; /**< event structure for recvs */
|
2005-07-04 02:45:48 +04:00
|
|
|
opal_mutex_t tcp_lock; /**< lock for accessing module state */
|
2005-07-03 20:22:16 +04:00
|
|
|
opal_list_t tcp_events; /**< list of pending events (accepts) */
|
|
|
|
opal_list_t tcp_msg_post; /**< list of recieves user has posted */
|
|
|
|
opal_list_t tcp_msg_recv; /**< list of recieved messages */
|
2005-10-25 17:48:08 +04:00
|
|
|
opal_list_t tcp_msg_completed; /**< list of completed messages */
|
2005-07-04 02:45:48 +04:00
|
|
|
opal_mutex_t tcp_match_lock; /**< lock held while searching/posting messages */
|
|
|
|
opal_condition_t tcp_match_cond; /**< condition variable used in finalize */
|
2004-09-30 19:09:29 +04:00
|
|
|
int tcp_match_count; /**< number of matched recvs in progress */
|
2004-09-02 03:07:40 +04:00
|
|
|
int tcp_debug; /**< debug level */
|
2004-07-13 02:46:57 +04:00
|
|
|
};
|
2004-08-16 23:39:54 +04:00
|
|
|
|
2004-08-10 03:07:53 +04:00
|
|
|
/**
|
|
|
|
* Convenience Typedef
|
|
|
|
*/
|
2004-07-13 02:46:57 +04:00
|
|
|
typedef struct mca_oob_tcp_component_t mca_oob_tcp_component_t;
|
|
|
|
|
2004-10-22 20:06:05 +04:00
|
|
|
OMPI_COMP_EXPORT extern mca_oob_tcp_component_t mca_oob_tcp_component;
|
2004-07-13 02:46:57 +04:00
|
|
|
|
|
|
|
|
2004-07-01 18:49:54 +04:00
|
|
|
#if defined(c_plusplus) || defined(__cplusplus)
|
|
|
|
}
|
|
|
|
#endif
|
|
|
|
|
|
|
|
#endif /* MCA_OOB_TCP_H_ */
|
|
|
|
|