@inproceedings{730,
  abstract     = {{Allocating resources to virtualized network functions and services to meet service level agreements is a challenging task for NFV management and orchestration systems. This becomes even more challenging when agile development methodologies, like DevOps, are applied. In such scenarios, management and orchestration systems are continuously facing new versions of functions and services which makes it hard to decide how much resources have to be allocated to them to provide the expected service performance. 
One solution for this problem is to support resource allocation decisions with performance behavior information obtained by profiling techniques applied to such network functions and services.

In this position paper, we analyze and discuss the components needed to generate such performance behavior information within the NFV DevOps workflow. We also outline research questions that identify open issues and missing pieces for a fully integrated NFV profiling solution. Further, we introduce a novel profiling mechanism that is able to profile virtualized network functions and entire network service chains under different resource constraints before they are deployed on production infrastructure.}},
  author       = {{Peuster, Manuel and Karl, Holger}},
  booktitle    = {{Fifth European Workshop on Software-Defined Networks, EWSDN 2016, Den Haag, The Netherlands, October 10-11, 2016}},
  location     = {{Den Haag}},
  pages        = {{7----12}},
  title        = {{{Understand Your Chains: Towards Performance Profile-Based Network Service Management}}},
  doi          = {{10.1109/EWSDN.2016.9}},
  year         = {{2016}},
}

@inproceedings{732,
  abstract     = {{Elastic deployments of virtualized network functions~(VNF) can automatically scale the amount of used resources in relation to their workload. This is often done by starting new VNF instances or stopping old ones. A problem of these scale operations is that most network functions are stateful and their internal state is not automatically migrated when traffic is redistributed in the deployment. As a result, mechanisms are needed to exchange or migrate internal network function state between VNF instances.

This paper presents a state management framework that creates a logically distributed state store on top of elastically deployed virtual network functions. We also introduce a novel programming model that provides both a local and a global view of the state to each VNF instance. We discuss the integration of our framework into existing network function virtualization architectures and compare the performance of our prototype to a centralized and a distributed state store solution.}},
  author       = {{Peuster, Manuel and Karl, Holger}},
  booktitle    = {{IEEE NetSoft Conference and Workshops, NetSoft 2016, Seoul, South Korea, June 6-10, 2016}},
  location     = {{Seoul}},
  pages        = {{6----10}},
  title        = {{{E-State: Distributed state management in elastic network function deployments}}},
  doi          = {{10.1109/NETSOFT.2016.7502432}},
  year         = {{2016}},
}

@inproceedings{738,
  abstract     = {{Virtualized network services consisting of multiple individual network functions are already today deployed across multiple sites, so called multi-PoP (points of presence) environments. This allows to improve service performance by optimizing its placement in the network. But prototyping and testing of these complex distributed software systems becomes extremely challenging. The reason is that not only the network service as such has to be tested but also its integration with management and orchestration systems. Existing solutions, like simulators, basic network emulators, or local cloud testbeds, do not support all aspects of these tasks.

To this end, we introduce MeDICINE, a novel NFV prototyping platform that is able to execute production-ready network functions, provided as software containers, in an emulated multi-PoP environment. These network functions can be controlled by any third-party management and orchestration system that connects to our platform through standard interfaces. Based on this, a developer can use our platform to prototype and test complex network services in a realistic environment running on his laptop.
}},
  author       = {{Peuster, Manuel and Karl, Holger and van Rossem, Steven}},
  booktitle    = {{IEEE Conference on Network Function Virtualization and Software Defined Networks (NFV-SDN)}},
  location     = {{Palo Alto}},
  title        = {{{MeDICINE: Rapid Prototyping of Production-Ready Network Services in Multi-PoP Environments}}},
  doi          = {{10.1109/NFV-SDN.2016.7919490}},
  year         = {{2016}},
}

@inproceedings{5595,
  author       = {{Wagner, Gerit and Prester, Julian and Roche, Maria and Benlian, Alexander and Schryen, Guido}},
  booktitle    = {{International Conference on Information Systems}},
  title        = {{{Factors Affecting the Scientific Impact of Literature Reviews: A Scientometric Study}}},
  year         = {{2016}},
}

@article{5617,
  abstract     = {{CAPTCHAs are challenge-response tests that aim at preventing unwanted machines, including bots, from accessing web services while providing easy access for humans. Recent advances in artificial-intelligence based attacks show that the level of security provided by many state-of-the-art text-based CAPTCHAs is declining. At the same time, techniques for distorting and obscuring the text, which are used to maintain the level of security, make text-based CAPTCHAs diffcult to solve for humans, and thereby further degrade usability. The need for developing alternative types of CAPTCHAs which improve both, the current security and usability levels, has been emphasized by several researchers. With this study, we contribute to research through (1) the development of two new face recognition CAPTCHAs (Farett-Gender and Farett-Gender&Age), (2) the security analysis of both procedures, and (3) the provision of empirical evidence that one of the suggested CAPTCHAs (Farett-Gender) is similar to Google's reCAPTCHA and better than KCAPTCHA concerning effectiveness (error rates), superior to both regarding learnability and satisfaction but not effciency.}},
  author       = {{Schryen, Guido and Wagner, Gerit and Schlegel, Alexander}},
  journal      = {{Computers & Security}},
  keywords     = {{CAPTCHA, Usability, Facial features, Gender classiffcation, Age classification, Face recognition reverse Turing test}},
  number       = {{July}},
  pages        = {{95--116}},
  publisher    = {{Elsevier}},
  title        = {{{Development of two novel face-recognition CAPTCHAs: a security and usability study}}},
  volume       = {{60}},
  year         = {{2016}},
}

@inproceedings{6739,
  author       = {{Wolters, Dennis and Gerth, Christian and Engels, Gregor}},
  booktitle    = {{Proceedings of the CAiSE'18 Forum at the 28th International Conference on Advanced Information Systems Engineering (CAiSE'16)}},
  pages        = {{89--96}},
  publisher    = {{CEUR-WS.org}},
  title        = {{{Modeling Cross-Device Systems with Use Case Diagrams}}},
  volume       = {{1612}},
  year         = {{2016}},
}

@inproceedings{166,
  abstract     = {{Network function virtualization and software-defined networking allow services consisting of virtual network functions to be designed and implemented with great flexibility by facilitating automatic deployments, migrations, and reconfigurations for services and their components. For extended flexibility, we go beyond seeing services as a fixed chain of functions. We present a YANG model for describing the service structure in deployment requests in a flexible way that enables changing the order of functions in case the order of traversing them does not affect the functionality of the service. Upon receiving such requests, the network orchestration system can choose the optimal composition of service components that gives the best results for placement of services in the network. This introduces new complexities to the placement problem by greatly increasing the number of possible ways a service can be composed. In this paper, we describe a heuristic solution that selects a Pareto set of the possible compositions of a service as well as possible combinations of different services, with respect to different resource requirements of the services. Our evaluations show that the selected combinations consist of representative samples of possible structures and requirements and therefore, can result in optimal or close-to-optimal placement results.}},
  author       = {{Dräxler, Sevil and Karl, Holger}},
  booktitle    = {{Proceedings of the 2nd International IEEE Conference on Network Softwarization (NetSoft)}},
  pages        = {{184----192}},
  title        = {{{Placement of Services with Flexible Structures Specified by a YANG Data Model}}},
  doi          = {{10.1109/NETSOFT.2016.7502412}},
  year         = {{2016}},
}

@inproceedings{11738,
  abstract     = {{In this contribution we investigate a priori signal-to-noise ratio (SNR) estimation, a crucial component of a single-channel speech enhancement system based on spectral subtraction. The majority of the state-of-the art a priori SNR estimators work in the power spectral domain, which is, however, not confirmed to be the optimal domain for the estimation. Motivated by the generalized spectral subtraction rule, we show how the estimation of the a priori SNR can be formulated in the so called generalized SNR domain. This formulation allows to generalize the widely used decision directed (DD) approach. An experimental investigation with different noise types reveals the superiority of the generalized DD approach over the conventional DD approach in terms of both the mean opinion score - listening quality objective measure and the output global SNR in the medium to high input SNR regime, while we show that the power spectrum is the optimal domain for low SNR. We further develop a parameterization which adjusts the domain of estimation automatically according to the estimated input global SNR. Index Terms: single-channel speech enhancement, a priori SNR estimation, generalized spectral subtraction}},
  author       = {{Chinaev, Aleksej and Haeb-Umbach, Reinhold}},
  booktitle    = {{INTERSPEECH 2016, San Francisco, USA}},
  title        = {{{A Priori SNR Estimation Using a Generalized Decision Directed Approach}}},
  year         = {{2016}},
}

@inproceedings{11743,
  abstract     = {{This contribution introduces a novel causal a priori signal-to-noise ratio (SNR) estimator for single-channel speech enhancement. To exploit the advantages of the generalized spectral subtraction, a normalized ?-order magnitude (NAOM) domain is introduced where an a priori SNR estimation is carried out. In this domain, the NAOM coefficients of noise and clean speech signals are modeled by a Weibull distribution and aWeibullmixturemodel (WMM), respectively. While the parameters of the noise model are calculated from the noise power spectral density estimates, the speechWMM parameters are estimated from the noisy signal by applying a causal Expectation-Maximization algorithm. Further a maximum a posteriori estimate of the a priori SNR is developed. The experiments in different noisy environments show the superiority of the proposed estimator compared to the well-known decision-directed approach in terms of estimation error, estimator variance and speech quality of the enhanced signals when used for speech enhancement.}},
  author       = {{Chinaev, Aleksej and Heitkaemper, Jens and Haeb-Umbach, Reinhold}},
  booktitle    = {{12. ITG Fachtagung Sprachkommunikation (ITG 2016)}},
  title        = {{{A Priori SNR Estimation Using Weibull Mixture Model}}},
  year         = {{2016}},
}

@inproceedings{11744,
  abstract     = {{A noise power spectral density (PSD) estimation is an indispensable component of speech spectral enhancement systems. In this paper we present a noise PSD tracking algorithm, which employs a noise presence probability estimate delivered by a deep neural network (DNN). The algorithm provides a causal noise PSD estimate and can thus be used in speech enhancement systems for communication purposes. An extensive performance comparison has been carried out with ten causal state-of-the-art noise tracking algorithms taken from the literature and categorized acc. to applied techniques. The experiments showed that the proposed DNN-based noise PSD tracker outperforms all competing methods with respect to all tested performance measures, which include the noise tracking performance and the performance of a speech enhancement system employing the noise tracking component.}},
  author       = {{Chinaev, Aleksej and Heymann, Jahn and Drude, Lukas and Haeb-Umbach, Reinhold}},
  booktitle    = {{12. ITG Fachtagung Sprachkommunikation (ITG 2016)}},
  title        = {{{Noise-Presence-Probability-Based Noise PSD Estimation by Using DNNs}}},
  year         = {{2016}},
}

@inproceedings{11751,
  author       = {{Drude, Lukas and Boeddeker, Christoph and Haeb-Umbach, Reinhold}},
  booktitle    = {{Proc. IEEE Intl. Conf. on Acoustics, Speech and Signal Processing (ICASSP)}},
  title        = {{{Blind Speech Separation based on Complex Spherical k-Mode Clustering}}},
  year         = {{2016}},
}

@inproceedings{11756,
  abstract     = {{Although complex-valued neural networks (CVNNs) â?? networks which can operate with complex arithmetic â?? have been around for a while, they have not been given reconsideration since the breakthrough of deep network architectures. This paper presents a critical assessment whether the novel tool set of deep neural networks (DNNs) should be extended to complex-valued arithmetic. Indeed, with DNNs making inroads in speech enhancement tasks, the use of complex-valued input data, specifically the short-time Fourier transform coefficients, is an obvious consideration. In particular when it comes to performing tasks that heavily rely on phase information, such as acoustic beamforming, complex-valued algorithms are omnipresent. In this contribution we recapitulate backpropagation in CVNNs, develop complex-valued network elements, such as the split-rectified non-linearity, and compare real- and complex-valued networks on a beamforming task. We find that CVNNs hardly provide a performance gain and conclude that the effort of developing the complex-valued counterparts of the building blocks of modern deep or recurrent neural networks can hardly be justified.}},
  author       = {{Drude, Lukas and Raj, Bhiksha and Haeb-Umbach, Reinhold}},
  booktitle    = {{INTERSPEECH 2016, San Francisco, USA}},
  title        = {{{On the appropriateness of complex-valued neural networks for speech enhancement}}},
  year         = {{2016}},
}

@inproceedings{11771,
  abstract     = {{This paper is concerned with speech presence probability estimation employing an explicit model of the temporal and spectral correlations of speech. An undirected graphical model is introduced, based on a Factor Graph formulation. It is shown that this undirected model cures some of the theoretical issues of an earlier directed graphical model. Furthermore, we formulate a message passing inference scheme based on an approximate graph factorization, identify this inference scheme as a particular message passing schedule based on the turbo principle and suggest further alternative schedules. The experiments show an improved performance over speech presence probability estimation based on an IID assumption, and a slightly better performance of the turbo schedule over the alternatives.}},
  author       = {{Glarner, Thomas and Mahdi Momenzadeh, Mohammad and Drude, Lukas and Haeb-Umbach, Reinhold}},
  booktitle    = {{12. ITG Fachtagung Sprachkommunikation (ITG 2016)}},
  title        = {{{Factor Graph Decoding for Speech Presence Probability Estimation}}},
  year         = {{2016}},
}

@inproceedings{11812,
  author       = {{Heymann, Jahn and Drude, Lukas and Haeb-Umbach, Reinhold}},
  booktitle    = {{Proc. IEEE Intl. Conf. on Acoustics, Speech and Signal Processing (ICASSP)}},
  title        = {{{Neural Network Based Spectral Mask Estimation for Acoustic Beamforming}}},
  year         = {{2016}},
}

@inproceedings{11829,
  abstract     = {{This contribution investigates Direction of Arrival (DoA) estimation using linearly arranged microphone arrays. We are going to develop a model for the DoA estimation error in a reverberant scenario and show the existence of a bias, that is a consequence of the linear arrangement and limited field of view (FoV) bias: First, the limited FoV leading to a clipping of the measurements, and, second, the angular distribution of the signal energy of the reflections being non-uniform. Since both issues are a consequence of the linear arrangement of the sensors, the bias arises largely independent of the kind of DoA estimator. The experimental evaluation demonstrates the existence of the bias for a selected number of DoA estimation methods and proves that the prediction from the developed theoretical model matches the simulation results.}},
  author       = {{Jacob, Florian and Haeb-Umbach, Reinhold}},
  booktitle    = {{12. ITG Fachtagung Sprachkommunikation (ITG 2016)}},
  title        = {{{On the Bias of Direction of Arrival Estimation Using Linear Microphone Arrays}}},
  year         = {{2016}},
}

@inproceedings{11834,
  abstract     = {{We present a system for the 4th CHiME challenge which significantly increases the performance for all three tracks with respect to the provided baseline system. The front-end uses a bi-directional Long Short-Term Memory (BLSTM)-based neural network to estimate signal statistics. These then steer a Generalized Eigenvalue beamformer. The back-end consists of a 22 layer deep Wide Residual Network and two extra BLSTM layers. Working on a whole utterance instead of frames allows us to refine Batch-Normalization. We also train our own BLSTM-based language model. Adding a discriminative speaker adaptation leads to further gains. The final system achieves a word error rate on the six channel real test data of 3.48%. For the two channel track we achieve 5.96% and for the one channel track 9.34%. This is the best reported performance on the challenge achieved by a single system, i.e., a configuration, which does not combine multiple systems. At the same time, our system is independent of the microphone configuration. We can thus use the same components for all three tracks.}},
  author       = {{Heymann, Jahn and Drude, Lukas and Haeb-Umbach, Reinhold}},
  booktitle    = {{Computer Speech and Language}},
  title        = {{{Wide Residual BLSTM Network with Discriminative Speaker Adaptation for Robust Speech Recognition}}},
  year         = {{2016}},
}

@article{11840,
  author       = {{Kinoshita, Keisuke and Delcroix, Marc and Gannot, Sharon and Habets, Emanuel A. P. and Haeb-Umbach, Reinhold and Kellermann, Walter and Leutnant, Volker and Maas, Roland and Nakatani, Tomohiro and Raj, Bhiksha and Sehr, Armin and Yoshioka, Takuya}},
  journal      = {{EURASIP Journal on Advances in Signal Processing}},
  title        = {{{A summary of the REVERB challenge: state-of-the-art and remaining challenges in reverberant speech processing research}}},
  year         = {{2016}},
}

@inproceedings{11908,
  abstract     = {{This paper describes automatic speech recognition (ASR) systems developed jointly by RWTH, UPB and FORTH for the 1ch, 2ch and 6ch track of the 4th CHiME Challenge. In the 2ch and 6ch tracks the final system output is obtained by a Confusion Network Combination (CNC) of multiple systems. The Acoustic Model (AM) is a deep neural network based on Bidirectional Long Short-Term Memory (BLSTM) units. The systems differ by front ends and training sets used for the acoustic training. The model for the 1ch track is trained without any preprocessing. For each front end we trained and evaluated individual acoustic models. We compare the ASR performance of different beamforming approaches: a conventional superdirective beamformer [1] and an MVDR beamformer as in [2], where the steering vector is estimated based on [3]. Furthermore we evaluated a BLSTM supported Generalized Eigenvalue beamformer using NN-GEV [4]. The back end is implemented using RWTH?s open-source toolkits RASR [5], RETURNN [6] and rwthlm [7]. We rescore lattices with a Long Short-Term Memory (LSTM) based language model. The overall best results are obtained by a system combination that includes the lattices from the system of UPB?s submission [8]. Our final submission scored second in each of the three tracks of the 4th CHiME Challenge.}},
  author       = {{Menne, Tobias and Heymann, Jahn and Alexandridis, Anastasios and Irie, Kazuki and Zeyer, Albert and Kitza, Markus and Golik, Pavel and Kulikov, Ilia and Drude, Lukas and Schlüter, Ralf and Ney, Hermann and Haeb-Umbach, Reinhold and Mouchtaris, Athanasios}},
  booktitle    = {{Computer Speech and Language}},
  title        = {{{The RWTH/UPB/FORTH System Combination for the 4th CHiME Challenge Evaluation}}},
  year         = {{2016}},
}

@inproceedings{11920,
  abstract     = {{In this paper we demonstrate an algorithm to learn words from speech using non-parametric Bayesian hierarchical models in an unsupervised setting. We exploit the assumption of a hierarchical structure of speech, namely the formation of spoken words as a sequence of phonemes. We employ the Nested Hierarchical Pitman-Yor Language Model, which allows an a priori unknown and possibly unlimited number of words. We assume the n-gram probabilities of words, the m-gram probabilities of phoneme sequences in words and the phoneme sequences of the words themselves as latent variables to be learned. We evaluate the algorithm on a cross language task using an existing speech recognizer trained on English speech to decode speech in the Xitsonga language supplied for the 2015 ZeroSpeech challenge. We apply the learning algorithm on the resulting phoneme graphs and achieve the highest token precision and F score compared to present systems.}},
  author       = {{Walter, Oliver and Haeb-Umbach, Reinhold}},
  booktitle    = {{38th German Conference on Pattern Recognition (GCPR 2016)}},
  title        = {{{Unsupervised Word Discovery from Speech using Bayesian Hierarchical Models}}},
  year         = {{2016}},
}

@inbook{12935,
  author       = {{Komprecht, Anna Maria and Röwenstrunk, Daniel}},
  booktitle    = {{„Ei, dem alten Herrn zoll’ ich Achtung gern“}},
  publisher    = {{Allitera Verlag, München}},
  title        = {{{Projektmanagement in digitalen Forschungsprojekten}}},
  doi          = {{10.25366/2018.33}},
  year         = {{2016}},
}

